@misc{indiciae18aff5561126, title = {Where Visual Speech Meets Language: VSP-LLM Framework for Efficient and Context-Aware Visual Speech Processing}, author = {Jeong Hun Yeo and Seunghee Han and Minsu Kim and Yong Man Ro}, year = {2024}, url = {https://arxiv.org/abs/2402.15151}, note = {Source identifier: 2402.15151} }