@misc{indiciaeb9d54759f379, title = {Reinforcement Learning vs. Distillation: Understanding Accuracy and Capability in LLM Reasoning}, author = {Minwu Kim and Anubhav Shrestha and Safal Shrestha and Aadim Nepal and Keith Ross}, year = {2025}, url = {https://arxiv.org/abs/2505.14216}, note = {Source identifier: 2505.14216} }