@misc{indiciae2a5dd55f9b40, title = {Cascade-CLIP: Cascaded Vision-Language Embeddings Alignment for Zero-Shot Semantic Segmentation}, author = {Yunheng Li and ZhongYu Li and Quansheng Zeng and Qibin Hou and Ming-Ming Cheng}, year = {2024}, url = {https://arxiv.org/abs/2406.00670}, note = {Source identifier: 2406.00670} }