@misc{indiciaeb74aba6cb2c0, title = {Antidote: Post-fine-tuning Safety Alignment for Large Language Models against Harmful Fine-tuning}, author = {Tiansheng Huang and Gautam Bhattacharya and Pratik Joshi and Josh Kimball and Ling Liu}, year = {2025}, url = {https://arxiv.org/abs/2408.09600}, note = {Source identifier: 2408.09600} }