@misc{indiciae627b41769ce9, title = {SyncVSR: Data-Efficient Visual Speech Recognition with End-to-End Crossmodal Audio Token Synchronization}, author = {Young Jin Ahn and Jungwoo Park and Sangha Park and Jonghyun Choi and Kee-Eung Kim}, year = {2024}, url = {https://arxiv.org/abs/2406.12233}, note = {Source identifier: 2406.12233} }