@misc{indiciaebb6e4544e3c3, title = {GAPO: Learning Preferential Prompt through Generative Adversarial Policy Optimization}, author = {Zhouhong Gu and Xingzhou Chen and Xiaoran Shi and Tao Wang and Suhang Zheng and Tianyu Li and Hongwei Feng and Yanghua Xiao}, year = {2025}, url = {https://arxiv.org/abs/2503.20194}, note = {Source identifier: 2503.20194} }