<?xml version='1.0' encoding='utf-8'?>
<?xml-stylesheet type="text/xsl" href="/v2/static/oai2.xsl"?>
<OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd">
  <responseDate>2026-10-08T21:22:53Z</responseDate>
  <request identifier="oai:figshare.com:article/34068783" metadataPrefix="oai_dc" verb="GetRecord">https://api.figshare.com/v2/oai</request>
  <GetRecord>
    <record>
      <header>
        <identifier>oai:figshare.com:article/34068783</identifier>
        <datestamp>2026-10-05T09:26:45Z</datestamp>
        <setSpec>category_29200</setSpec>
        <setSpec>category_29203</setSpec>
        <setSpec>item_type_3</setSpec>
        <setSpec>month_year_10_2026</setSpec>
      </header>
      <metadata>
        <oai_dc:dc xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance"  xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
          <dc:title>Replication package for "When Does LLM Test Repair Destroy Bug-Revealing Tests? A Cost-Accounted Study on Real Repository Defects"</dc:title>
          <dc:creator>Zehao Li (25144737)</dc:creator>
          <dc:subject>Software testing, verification and validation</dc:subject>
          <dc:subject>Software engineering not elsewhere classified</dc:subject>
          <dc:subject>software testing</dc:subject>
          <dc:subject>test repair</dc:subject>
          <dc:subject>test oracle</dc:subject>
          <dc:subject>large language models</dc:subject>
          <dc:subject>bug reproduction</dc:subject>
          <dc:subject>SWT-bench</dc:subject>
          <dc:subject>empirical study</dc:subject>
          <dc:description>&lt;p dir="ltr"&gt;Replication package for an empirical study of how iterative LLM test repair destroys verified bug-revealing (fail-to-pass) tests on real repository defects (SWT-bench Verified, 229 instances, three LLMs). It contains the study protocol with its dated amendments, the Docker-free execution harness, every LLM request and response with provider-reported cost, every recorded test version with its outcomes on the buggy and the fixed program version, the test-execution logs, the model-based audit samples, the official-harness cross-check, and the scripts that regenerate every table, figure and in-text number of the paper. Benchmark metadata are not redistributed: src/fetch_data.py downloads them from Hugging Face, and work/data_checksums.json verifies them. Code: MIT. Data: CC BY 4.0.&lt;/p&gt;</dc:description>
          <dc:date>2026-10-05T09:26:45Z</dc:date>
          <dc:type>Dataset</dc:type>
          <dc:type>Dataset</dc:type>
          <dc:identifier>10.6084/m9.figshare.34068783.v1</dc:identifier>
          <dc:relation>https://figshare.com/articles/dataset/Replication_package_for_When_Does_LLM_Test_Repair_Destroy_Bug-Revealing_Tests_A_Cost-Accounted_Study_on_Real_Repository_Defects_/34068783</dc:relation>
          <dc:rights>CC BY 4.0</dc:rights>
        </oai_dc:dc>
      </metadata>
    </record>
  </GetRecord>
</OAI-PMH>
