@misc{indiciaeb0ac47fb0342, title = {Linear Separability of Activation Representations after Supervised Fine-Tuning on Incorrect Responses: A Study of Synthetic Dishonesty in Large Language Models}, author = {Vahideh Zolfaghari}, year = {2026}, url = {https://arxiv.org/abs/2605.30381}, note = {Source identifier: 2605.30381} }