@misc{indiciaea463deb32ab8, title = {VIDiff: Translating Videos via Multi-Modal Instructions with Diffusion Models}, author = {Zhen Xing and Shuyuan Tu and Qi Dai and Zihao Zhang and Hui Zhang and Han Hu and Zuxuan Wu and Yu-Gang Jiang}, year = {2026}, url = {https://arxiv.org/abs/2311.18837}, note = {Source identifier: 2311.18837} }