@misc{indiciae01ef43d0fbcd, title = {Efficient Training for Multilingual Visual Speech Recognition: Pre-training with Discretized Visual Speech Representation}, author = {Minsu Kim and Jeong Hun Yeo and Se Jin Park and Hyeongseop Rha and Yong Man Ro}, year = {2024}, url = {https://arxiv.org/abs/2401.09802}, note = {Source identifier: 2401.09802} }