@misc{indiciaef483a7c38750, title = {To Switch or Not to Switch? Balanced Policy Switching in Offline Reinforcement Learning}, author = {Tao Ma and Xuzhi Yang and Zoltan Szabo}, year = {2025}, url = {https://arxiv.org/abs/2407.01837}, note = {Source identifier: 2407.01837} }