@misc{indiciae1a1e1c7ea98b, title = {Reward-agnostic Fine-tuning: Provable Statistical Benefits of Hybrid Reinforcement Learning}, author = {Gen Li and Wenhao Zhan and Jason D. Lee and Yuejie Chi and Yuxin Chen}, year = {2023}, url = {https://arxiv.org/abs/2305.10282}, note = {Source identifier: 2305.10282} }