@misc{indiciaee384e4509e54, title = {Integrating Fine-Grained Audio-Visual Evidence for Robust Multimodal Emotion Reasoning}, author = {Zhixian Zhao and Wenjie Tian and Lei Xie}, year = {2026}, url = {https://arxiv.org/abs/2601.18321}, note = {Source identifier: 2601.18321} }