@misc{indiciaeafabeedc3a94, title = {Inadequacies of Large Language Model Benchmarks in the Era of Generative Artificial Intelligence}, author = {Timothy R. McIntosh and Teo Susnjak and Nalin Arachchilage and Tong Liu and Paul Watters and Malka N. Halgamuge}, year = {2024}, doi = {10.1109/tai.2025.3569516}, url = {https://arxiv.org/abs/2402.09880}, note = {Source identifier: 2402.09880} }