@misc{indiciae6682d63aceca, title = {Baseline Defenses for Adversarial Attacks Against Aligned Language Models}, author = {Neel Jain and Avi Schwarzschild and Yuxin Wen and Gowthami Somepalli and John Kirchenbauer and Ping-yeh Chiang and Micah Goldblum and Aniruddha Saha and Jonas Geiping and Tom Goldstein}, year = {2023}, url = {https://arxiv.org/abs/2309.00614}, note = {Source identifier: 2309.00614} }