@misc{indiciaefe32303b2144, title = {Reinforcement Learning for Chain of Thought Compression with One-Domain-to-All Generalization}, author = {Hanyu Li and Jiangshan Duo and Bofei Gao and Hailin Zhang and Sujian Li and Xiaotie Deng and Liang Zhao}, year = {2026}, url = {https://arxiv.org/abs/2601.06052}, note = {Source identifier: 2601.06052} }