@misc{indiciaebe4dbc4c9ff5, title = {Reinforcement Learning Upside Down: Don't Predict Rewards -- Just Map Them to Actions}, author = {Juergen Schmidhuber}, year = {2020}, url = {https://arxiv.org/abs/1912.02875}, note = {Source identifier: 1912.02875} }