@misc{indiciaefd56c3068f00, title = {Where-to-Learn: Analytical Policy Gradient Directed Exploration for On-Policy Robotic Reinforcement Learning}, author = {Leixin Chang and Xinchen Yao and Ben Liu and Liangjing Yang and Hua Chen}, year = {2026}, doi = {10.1109/lra.2026.3678143}, url = {https://arxiv.org/abs/2603.27317}, note = {Source identifier: 2603.27317} }