@misc{indiciaebc4da655bf9f, title = {Sampling Efficient Deep Reinforcement Learning through Preference-Guided Stochastic Exploration}, author = {Wenhui Huang and Cong Zhang and Jingda Wu and Xiangkun He and Jie Zhang and Chen Lv}, year = {2022}, url = {https://arxiv.org/abs/2206.09627}, note = {Source identifier: 2206.09627} }