@misc{indiciae941307a697e9, title = {TADPO: Reinforcement Learning Goes Off-road}, author = {Zhouchonghao Wu and Raymond Song and Vedant Mundheda and Luis E. Navarro-Serment and Christof Schoenborn and Jeff Schneider}, year = {2026}, url = {https://arxiv.org/abs/2603.05995}, note = {Source identifier: 2603.05995} }