@misc{indiciaee734e71b1a46, title = {Efficient Preference-based Reinforcement Learning via Aligned Experience Estimation}, author = {Fengshuo Bai and Rui Zhao and Hongming Zhang and Sijia Cui and Ying Wen and Yaodong Yang and Bo Xu and Lei Han}, year = {2024}, url = {https://arxiv.org/abs/2405.18688}, note = {Source identifier: 2405.18688} }