@misc{indiciae27a2519a3a4a, title = {Leveraging Large Language Models in Visual Speech Recognition: Model Scaling, Context-Aware Decoding, and Iterative Polishing}, author = {Zehua Liu and Xiaolou Li and Li Guo and Lantian Li and Dong Wang}, year = {2025}, url = {https://arxiv.org/abs/2506.02012}, note = {Source identifier: 2506.02012} }