@misc{indiciae8994a4461149, title = {The Multimodal Information Based Speech Processing (MISP) 2025 Challenge: Audio-Visual Diarization and Recognition}, author = {Ming Gao and Shilong Wu and Hang Chen and Jun Du and Chin-Hui Lee and Shinji Watanabe and Jingdong Chen and Siniscalchi Sabato Marco and Odette Scharenborg}, year = {2025}, url = {https://arxiv.org/abs/2505.13971}, note = {Source identifier: 2505.13971} }