Syed Fakhar Abbas Naqvi, Hina Anwar. Benchmarking Performance and Sustainability Trade-Offs in LLMs for Python Bug-Fixing Tasks (Software, Source Code). Schloss Dagstuhl – Leibniz-Zentrum für Informatik (2026)
@misc{dagstuhl-artifact-28073,
title = {{Benchmarking Performance and Sustainability Trade-Offs in LLMs for Python Bug-Fixing Tasks}},
author = {Naqvi, Syed Fakhar Abbas and Anwar, Hina},
note = {Software, swhId: \href{https://archive.softwareheritage.org/swh:1:dir:e2f5ea38e5e8b633db8ef0baf818d58b194468a9;origin=https://github.com/syedfakhar25/empirical-llm-bugfix-benchmark;visit=swh:1:snp:30dfb758d4b717c0deb55403872688c9af37f075;anchor=swh:1:rev:e3cc68f625c69eacd334437f1c9c1343aae6b348}{\texttt{swh:1:dir:e2f5ea38e5e8b633db8ef0baf818d58b194468a9}} (visited on 2026-10-05)},
url = {https://github.com/syedfakhar25/empirical-llm-bugfix-benchmark/tree/main},
doi = {10.4230/artifacts.28073},
}
Published in: LIPIcs, Volume 394, 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)
Feray Gulu-zada and Hina Anwar. Balancing Green and Clean Code: Prompting for LLM-Based Refactoring. In 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026). Leibniz International Proceedings in Informatics (LIPIcs), Volume 394, pp. 29:1-29:20, Schloss Dagstuhl – Leibniz-Zentrum für Informatik (2026)
@InProceedings{guluzada_et_al:LIPIcs.ESEM.2026.29,
author = {Gulu-zada, Feray and Anwar, Hina},
title = {{Balancing Green and Clean Code: Prompting for LLM-Based Refactoring}},
booktitle = {20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)},
pages = {29:1--29:20},
series = {Leibniz International Proceedings in Informatics (LIPIcs)},
ISBN = {978-3-95977-450-5},
ISSN = {1868-8969},
year = {2026},
volume = {394},
editor = {Feldt, Robert and Paasivaara, Maria and Mendez, Daniel and Wagner, Stefan and Bar\'{o}n, Marvin Mu\~{n}oz},
publisher = {Schloss Dagstuhl -- Leibniz-Zentrum f{\"u}r Informatik},
address = {Dagstuhl, Germany},
URL = {https://drops.dagstuhl.de/entities/document/10.4230/LIPIcs.ESEM.2026.29},
URN = {urn:nbn:de:0030-drops-279971},
doi = {10.4230/LIPIcs.ESEM.2026.29},
annote = {Keywords: Large language models, code refactoring, energy efficiency, maintainability, green software, prompt engineering}
}
Published in: LIPIcs, Volume 394, 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)
Syed Fakhar Abbas Naqvi and Hina Anwar. Bigger Is Not Always Better: Performance and Sustainability Trade-Offs in LLMs for Python Bug-Fixing Tasks. In 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026). Leibniz International Proceedings in Informatics (LIPIcs), Volume 394, pp. 34:1-34:20, Schloss Dagstuhl – Leibniz-Zentrum für Informatik (2026)
@InProceedings{naqvi_et_al:LIPIcs.ESEM.2026.34,
author = {Naqvi, Syed Fakhar Abbas and Anwar, Hina},
title = {{Bigger Is Not Always Better: Performance and Sustainability Trade-Offs in LLMs for Python Bug-Fixing Tasks}},
booktitle = {20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)},
pages = {34:1--34:20},
series = {Leibniz International Proceedings in Informatics (LIPIcs)},
ISBN = {978-3-95977-450-5},
ISSN = {1868-8969},
year = {2026},
volume = {394},
editor = {Feldt, Robert and Paasivaara, Maria and Mendez, Daniel and Wagner, Stefan and Bar\'{o}n, Marvin Mu\~{n}oz},
publisher = {Schloss Dagstuhl -- Leibniz-Zentrum f{\"u}r Informatik},
address = {Dagstuhl, Germany},
URL = {https://drops.dagstuhl.de/entities/document/10.4230/LIPIcs.ESEM.2026.34},
URN = {urn:nbn:de:0030-drops-280023},
doi = {10.4230/LIPIcs.ESEM.2026.34},
annote = {Keywords: Large Language Models, Bug Fixing, Energy Consumption, Empirical Software Engineering, Sustainability}
}