@misc{indiciae431577a7ca36, title = {Self-play with Execution Feedback: Improving Instruction-following Capabilities of Large Language Models}, author = {Guanting Dong and Keming Lu and Chengpeng Li and Tingyu Xia and Bowen Yu and Chang Zhou and Jingren Zhou}, year = {2024}, url = {https://arxiv.org/abs/2406.13542}, note = {Source identifier: 2406.13542} }