@misc{indiciaec12dec520a00, title = {DR3: Value-Based Deep Reinforcement Learning Requires Explicit Regularization}, author = {Aviral Kumar and Rishabh Agarwal and Tengyu Ma and Aaron Courville and George Tucker and Sergey Levine}, year = {2021}, url = {https://arxiv.org/abs/2112.04716}, note = {Source identifier: 2112.04716} }