@misc{indiciaee989b42ddfdd, title = {What Makes VLMs Robust? Towards Reconciling Robustness and Accuracy in Vision-Language Models}, author = {Sen Nie and Jie Zhang and Zhongqi Wang and Zhaoyang Wei and Shiguang Shan and Xilin Chen}, year = {2026}, url = {https://arxiv.org/abs/2603.12799}, note = {Source identifier: 2603.12799} }