@misc{indiciae40b2e9234618, title = {VMonarch: Efficient Video Diffusion Transformers with Structured Attention}, author = {Cheng Liang and Haoxian Chen and Liang Hou and Qi Fan and Gangshan Wu and Xin Tao and Limin Wang}, year = {2026}, url = {https://arxiv.org/abs/2601.22275}, note = {Source identifier: 2601.22275} }