@misc{indiciae32f8178d06e7, title = {Vulnerability Mitigation for Safety-Aligned Language Models via Debiasing}, author = {Thien Q. Tran and Akifumi Wachi and Rei Sato and Takumi Tanabe and Youhei Akimoto}, year = {2025}, url = {https://arxiv.org/abs/2502.02153}, note = {Source identifier: 2502.02153} }