@misc{indiciae71bac96233a7, title = {Successive Training Stages and Large Language Model Persuasion: Effects of Misalignment, Supervised Fine-Tuning, and Preference Optimization}, author = {Antony Dalmiere and Pascal Marchand and Guillaume Auriol and Vincent Nicomette}, year = {2026}, url = {https://arxiv.org/abs/2610.09964}, note = {Source identifier: 2610.09964} }