@misc{indiciaec88464e61dfa, title = {SpeakerSleuth: Can Large Audio-Language Models Judge Speaker Consistency across Multi-turn Dialogues?}, author = {Jonggeun Lee and Junseong Pyo and Gyuhyeon Seo and Yohan Jo}, year = {2026}, url = {https://arxiv.org/abs/2601.04029}, note = {Source identifier: 2601.04029} }