@misc{indiciae32b089d9846e, title = {Mimir: Improving Video Diffusion Models for Precise Text Understanding}, author = {Shuai Tan and Biao Gong and Yutong Feng and Kecheng Zheng and Dandan Zheng and Shuwei Shi and Yujun Shen and Jingdong Chen and Ming Yang}, year = {2024}, url = {https://arxiv.org/abs/2412.03085}, note = {Source identifier: 2412.03085} }