@misc{indiciae389efbcd2a60, title = {Efficient Audiovisual Speech Processing via MUTUD: Multimodal Training and Unimodal Deployment}, author = {Joanna Hong and Sanjeel Parekh and Honglie Chen and Jacob Donley and Ke Tan and Buye Xu and Anurag Kumar}, year = {2026}, url = {https://arxiv.org/abs/2501.18157}, note = {Source identifier: 2501.18157} }