@misc{indiciae7617ebfb87f6, title = {RL Is Neither a Panacea Nor a Mirage: Understanding Supervised vs. Reinforcement Learning Fine-Tuning for LLMs}, author = {Hangzhan Jin and Sicheng Lv and Sifan Wu and Mohammad Hamdaqa}, year = {2025}, url = {https://arxiv.org/abs/2508.16546}, note = {Source identifier: 2508.16546} }