@misc{indiciaef4317b2ba90d, title = {LPPG-RL: Lexicographically Projected Policy Gradient Reinforcement Learning with Subproblem Exploration}, author = {Ruiyu Qiu and Rui Wang and Guanghui Yang and Xiang Li and Zhijiang Shao}, year = {2025}, url = {https://arxiv.org/abs/2511.08339}, note = {Source identifier: 2511.08339} }