@misc{indiciaebbf2dbe43c50, title = {Decomposing and Interpreting Image Representations via Text in ViTs Beyond CLIP}, author = {Sriram Balasubramanian and Samyadeep Basu and Soheil Feizi}, year = {2024}, url = {https://arxiv.org/abs/2406.01583}, note = {Source identifier: 2406.01583} }