@misc{indiciae77acec64b6a4, title = {Extending CLIP's Image-Text Alignment to Referring Image Segmentation}, author = {Seoyeon Kim and Minguk Kang and Dongwon Kim and Jaesik Park and Suha Kwak}, year = {2024}, url = {https://arxiv.org/abs/2306.08498}, note = {Source identifier: 2306.08498} }