@misc{indiciae65179da5f643, title = {CLIP as RNN: Segment Countless Visual Concepts without Training Endeavor}, author = {Shuyang Sun and Runjia Li and Philip Torr and Xiuye Gu and Siyang Li}, year = {2024}, url = {https://arxiv.org/abs/2312.07661}, note = {Source identifier: 2312.07661} }