@misc{indiciaebfa4dec17a24, title = {COMET: Concept Space Dissection of the Modality Gap in Audio-Text Multimodal Contrastive Embeddings}, author = {Yonggang Zhu and Liting Gao and Aidong Men and Wenwu Wang}, year = {2026}, url = {https://arxiv.org/abs/2605.29628}, note = {Source identifier: 2605.29628} }