@misc{indiciaefac2c93f2a9f, title = {Multi-Grained Cross-modal Alignment for Learning Open-vocabulary Semantic Segmentation from Text Supervision}, author = {Yajie Liu and Pu Ge and Qingjie Liu and Di Huang}, year = {2024}, url = {https://arxiv.org/abs/2403.03707}, note = {Source identifier: 2403.03707} }