@misc{indiciae09af2fd881f3, title = {Narrowing Information Bottleneck Theory for Multimodal Image-Text Representations Interpretability}, author = {Zhiyu Zhu and Zhibo Jin and Jiayu Zhang and Nan Yang and Jiahao Huang and Jianlong Zhou and Fang Chen}, year = {2025}, url = {https://arxiv.org/abs/2502.14889}, note = {Source identifier: 2502.14889} }