@misc{indiciae10c3d805e26f, title = {In-Context Reinforcement Learning From Suboptimal Historical Data}, author = {Juncheng Dong and Moyang Guo and Ethan X. Fang and Zhuoran Yang and Vahid Tarokh}, year = {2026}, url = {https://arxiv.org/abs/2601.20116}, note = {Source identifier: 2601.20116} }