@misc{indiciaeb4fe8e9afb4c, title = {GVDIFF: Grounded Text-to-Video Generation with Diffusion Models}, author = {Huanzhang Dou and Ruixiang Li and Wei Su and Xi Li}, year = {2024}, url = {https://arxiv.org/abs/2407.01921}, note = {Source identifier: 2407.01921} }