@misc{indiciae2507a5eb6faf, title = {Reward Is Enough: LLMs Are In-Context Reinforcement Learners}, author = {Kefan Song and Amir Moeini and Peng Wang and Lei Gong and Rohan Chandra and Shangtong Zhang and Yanjun Qi}, year = {2026}, url = {https://arxiv.org/abs/2506.06303}, note = {Source identifier: 2506.06303} }