@misc{indiciaea5541ea904ba, title = {XLAVS-R: Cross-Lingual Audio-Visual Speech Representation Learning for Noise-Robust Speech Perception}, author = {HyoJung Han and Mohamed Anwar and Juan Pino and Wei-Ning Hsu and Marine Carpuat and Bowen Shi and Changhan Wang}, year = {2024}, url = {https://arxiv.org/abs/2403.14402}, note = {Source identifier: 2403.14402} }