@misc{indiciaeee1373516f82, title = {Reinforcement World Model Learning for LLM-based Agents}, author = {Xiao Yu and Baolin Peng and Ruize Xu and Yelong Shen and Pengcheng He and Suman Nath and Nikhil Singh and Jiangfeng Gao and Zhou Yu}, year = {2026}, url = {https://arxiv.org/abs/2602.05842}, note = {Source identifier: 2602.05842} }