@misc{indiciae9cf2a724a5a4, title = {Robot Policy Learning with Temporal Optimal Transport Reward}, author = {Yuwei Fu and Haichao Zhang and Di Wu and Wei Xu and Benoit Boulet}, year = {2024}, url = {https://arxiv.org/abs/2410.21795}, note = {Source identifier: 2410.21795} }