@misc{indiciaee8af1b67f947, title = {Traversing Pareto Optimal Policies: Provably Efficient Multi-Objective Reinforcement Learning}, author = {Shuang Qiu and Dake Zhang and Rui Yang and Boxiang Lyu and Tong Zhang}, year = {2024}, url = {https://arxiv.org/abs/2407.17466}, note = {Source identifier: 2407.17466} }