@misc{indiciae48ed42928d86, title = {LLMs Are In-Context Bandit Reinforcement Learners}, author = {Giovanni Monea and Antoine Bosselut and Kianté Brantley and Yoav Artzi}, year = {2025}, url = {https://arxiv.org/abs/2410.05362}, note = {Source identifier: 2410.05362} }