@misc{indiciae3a222b999d89, title = {DoubleHelix: Structured Cross-Modal Fusion for Audio-Visual Speech Recognition with LLMs}, author = {Ziwei Cheng and Zhenhua Tan and Zhuomin Zhu}, year = {2026}, url = {https://arxiv.org/abs/2607.29112}, note = {Source identifier: 2607.29112} }