@misc{indiciae851a7e403679, title = {AI Alignment through Reinforcement Learning from Human Feedback? Contradictions and Limitations}, author = {Adam Dahlgren Lindström and Leila Methnani and Lea Krause and Petter Ericson and Íñigo Martínez de Rituerto de Troya and Dimitri Coelho Mollo and Roel Dobbe}, year = {2024}, url = {https://arxiv.org/abs/2406.18346}, note = {Source identifier: 2406.18346} }