@misc{indiciae06c2ab589479, title = {Learning to Reason as Action Abstractions with Scalable Mid-Training RL}, author = {Shenao Zhang and Donghan Yu and Yihao Feng and Bowen Jin and Zhaoran Wang and John Peebles and Zirui Wang}, year = {2025}, url = {https://arxiv.org/abs/2509.25810}, note = {Source identifier: 2509.25810} }