@misc{indiciae98518591718f, title = {Show Me What and Tell Me How: Video Synthesis via Multimodal Conditioning}, author = {Ligong Han and Jian Ren and Hsin-Ying Lee and Francesco Barbieri and Kyle Olszewski and Shervin Minaee and Dimitris Metaxas and Sergey Tulyakov}, year = {2022}, url = {https://arxiv.org/abs/2203.02573}, note = {Source identifier: 2203.02573} }