@misc{indiciaef1a6df74ce09, title = {MoDiT: Learning Highly Consistent 3D Motion Coefficients with Diffusion Transformer for Talking Head Generation}, author = {Yucheng Wang and Dan Xu}, year = {2025}, url = {https://arxiv.org/abs/2507.05092}, note = {Source identifier: 2507.05092} }