@misc{indiciae178b4dc1469b, title = {Video Diffusion Transformers are In-Context Learners}, author = {Zhengcong Fei and Di Qiu and Debang Li and Changqian Yu and Mingyuan Fan}, year = {2025}, url = {https://arxiv.org/abs/2412.10783}, note = {Source identifier: 2412.10783} }