@misc{indiciae92ad82fa7b09, title = {Cross-Modal Global Interaction and Local Alignment for Audio-Visual Speech Recognition}, author = {Yuchen Hu and Ruizhe Li and Chen Chen and Heqing Zou and Qiushi Zhu and Eng Siong Chng}, year = {2023}, url = {https://arxiv.org/abs/2305.09212}, note = {Source identifier: 2305.09212} }