@misc{indiciaedab3a98cd84f, title = {Automatic Generation of Benchmarks and Reliable LLM Judgment for Code Tasks}, author = {Eitan Farchi and Shmulik Froimovich and Rami Katan and Orna Raz}, year = {2024}, url = {https://arxiv.org/abs/2410.21071}, note = {Source identifier: 2410.21071} }