@misc{indiciae81ce750340bc, title = {Prompted Policy Search: Reinforcement Learning through Linguistic and Numerical Reasoning in LLMs}, author = {Yifan Zhou and Sachin Grover and Mohamed El Mistiri and Kamalesh Kalirathnam and Pratyush Kerhalkar and Swaroop Mishra and Neelesh Kumar and Sanket Gaurav and Oya Aran and Heni Ben Amor}, year = {2025}, url = {https://arxiv.org/abs/2511.21928}, note = {Source identifier: 2511.21928} }