@misc{indiciae5b2fbe63ece3, title = {Consistent Human Image and Video Generation with Spatially Conditioned Diffusion}, author = {Mingdeng Cao and Chong Mou and Ziyang Yuan and Xintao Wang and Zhaoyang Zhang and Ying Shan and Yinqiang Zheng}, year = {2024}, url = {https://arxiv.org/abs/2412.14531}, note = {Source identifier: 2412.14531} }