@misc{indiciae31b45d9e4eca, title = {Images are Achilles' Heel of Alignment: Exploiting Visual Vulnerabilities for Jailbreaking Multimodal Large Language Models}, author = {Yifan Li and Hangyu Guo and Kun Zhou and Wayne Xin Zhao and Ji-Rong Wen}, year = {2025}, url = {https://arxiv.org/abs/2403.09792}, note = {Source identifier: 2403.09792} }