@misc{indiciaebc2c62621efe, title = {Active Advantage-Aligned Online Reinforcement Learning with Offline Data}, author = {Xuefeng Liu and Hung T. C. Le and Siyu Chen and Rick Stevens and Zhuoran Yang and Matthew R. Walter and Yuxin Chen}, year = {2026}, url = {https://arxiv.org/abs/2502.07937}, note = {Source identifier: 2502.07937} }