@misc{indiciaef2cb4078321e, title = {Pre-training LLM without Learning Rate Decay Enhances Supervised Fine-Tuning}, author = {Kazuki Yano and Shun Kiyono and Sosuke Kobayashi and Sho Takase and Jun Suzuki}, year = {2026}, url = {https://arxiv.org/abs/2603.16127}, note = {Source identifier: 2603.16127} }