@misc{indiciaeb5d568553182, title = {Reusing Embeddings: Reproducible Reward Model Research in Large Language Model Alignment without GPUs}, author = {Hao Sun and Yunyi Shen and Jean-Francois Ton and Mihaela van der Schaar}, year = {2025}, url = {https://arxiv.org/abs/2502.04357}, note = {Source identifier: 2502.04357} }