@misc{indiciaecd3f840b864d, title = {Long-Horizon Language Model Reinforcement Learning via Progressive Point Matching}, author = {Preston Fu and Kevin Frans and Oleh Rybkin and Sergey Levine and Aviral Kumar}, year = {2026}, url = {https://arxiv.org/abs/2609.07303}, note = {Source identifier: 2609.07303} }