@misc{indiciae905a5fa7417a, title = {Zero-AVSR: Zero-Shot Audio-Visual Speech Recognition with LLMs by Learning Language-Agnostic Speech Representations}, author = {Jeong Hun Yeo and Minsu Kim and Chae Won Kim and Stavros Petridis and Yong Man Ro}, year = {2025}, url = {https://arxiv.org/abs/2503.06273}, note = {Source identifier: 2503.06273} }