@misc{indiciaeb30ffbd12160, title = {Cosine Misleads: Auxiliary Losses Reshape Vision Language Models, Not Their Latents}, author = {XiuYu Zhang and Junfeng Fang and Zhenkai Liang}, year = {2026}, url = {https://arxiv.org/abs/2606.05753}, note = {Source identifier: 2606.05753} }