@misc{indiciaeacbf98e39603, title = {Long-Horizon Model-Based Offline Reinforcement Learning Without Explicit Conservatism}, author = {Tianwei Ni and Esther Derman and Vineet Jain and Vincent Taboga and Siamak Ravanbakhsh and Pierre-Luc Bacon}, year = {2026}, url = {https://arxiv.org/abs/2512.04341}, note = {Source identifier: 2512.04341} }