@misc{indiciae694f259783eb, title = {Robust CLIP: Unsupervised Adversarial Fine-Tuning of Vision Embeddings for Robust Large Vision-Language Models}, author = {Christian Schlarmann and Naman Deep Singh and Francesco Croce and Matthias Hein}, year = {2024}, url = {https://arxiv.org/abs/2402.12336}, note = {Source identifier: 2402.12336} }