@misc{indiciae920a58f2d3f9, title = {CogVideo: Large-scale Pretraining for Text-to-Video Generation via Transformers}, author = {Wenyi Hong and Ming Ding and Wendi Zheng and Xinghan Liu and Jie Tang}, year = {2022}, url = {https://arxiv.org/abs/2205.15868}, note = {Source identifier: 2205.15868} }