@misc{indiciae03241ea5bf65, title = {TAVID: Text-Driven Audio-Visual Interactive Dialogue Generation}, author = {Ji-Hoon Kim and Junseok Ahn and Doyeop Kwak and Joon Son Chung and Shinji Watanabe}, year = {2025}, url = {https://arxiv.org/abs/2512.20296}, note = {Source identifier: 2512.20296} }