@misc{indiciae2446e66805e8, title = {Leveraging Unimodal Self-Supervised Learning for Multimodal Audio-Visual Speech Recognition}, author = {Xichen Pan and Peiyu Chen and Yichen Gong and Helong Zhou and Xinbing Wang and Zhouhan Lin}, year = {2022}, url = {https://arxiv.org/abs/2203.07996}, note = {Source identifier: 2203.07996} }