@misc{indiciae0469606e4e9f, title = {Listen to the Latents: Self-Correcting Speech Recognition in Large Audio Language Models Through Hidden-State Interactions}, author = {Chan-Jan Hsu and Jaeyeon Kim and Chao-Han Huck Yang and Shinji Watanabe and Hung-yi Lee and Carlos Busso}, year = {2026}, url = {https://arxiv.org/abs/2609.02940}, note = {Source identifier: 2609.02940} }