@misc{indiciae39ddaaba11a3, title = {Safety Alignment of Large Language Models via Contrasting Safe and Harmful Distributions}, author = {Xiaoyun Zhang and Zhengyue Zhao and Wenxuan Shi and Kaidi Xu and Di Huang and Xing Hu}, year = {2025}, url = {https://arxiv.org/abs/2406.16743}, note = {Source identifier: 2406.16743} }