@misc{indiciae8e1a24e3a0f0, title = {FLAP: Fully-controllable Audio-driven Portrait Video Generation through 3D head conditioned diffusion model}, author = {Lingzhou Mu and Baiji Liu and Ruonan Zhang and Guiming Mo and Jiawei Jin and Kai Zhang and Haozhi Huang}, year = {2025}, url = {https://arxiv.org/abs/2502.19455}, note = {Source identifier: 2502.19455} }