<?xml version='1.0' encoding='utf-8'?>
<?xml-stylesheet type="text/xsl" href="/v2/static/oai2.xsl"?>
<OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd">
  <responseDate>2026-10-07T13:29:18Z</responseDate>
  <request identifier="oai:figshare.com:article/33972877" metadataPrefix="oai_dc" verb="GetRecord">https://api.figshare.com/v2/oai</request>
  <GetRecord>
    <record>
      <header>
        <identifier>oai:figshare.com:article/33972877</identifier>
        <datestamp>2026-09-23T11:52:46Z</datestamp>
        <setSpec>category_29200</setSpec>
        <setSpec>item_type_3</setSpec>
        <setSpec>month_year_09_2026</setSpec>
      </header>
      <metadata>
        <oai_dc:dc xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance"  xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
          <dc:title>Java dataset and results</dc:title>
          <dc:creator>peng (25096915)</dc:creator>
          <dc:subject>Software testing, verification and validation</dc:subject>
          <dc:subject>Unit test generation</dc:subject>
          <dc:subject>Agent skills</dc:subject>
          <dc:description>&lt;p dir="ltr"&gt;This dataset contains the Java benchmark projects and experimental results used to evaluate a skill-based framework for LLM-assisted Java unit-test generation and iterative test repair.&lt;/p&gt;&lt;p dir="ltr"&gt;The benchmark collection comprises Maven-based open-source Java projects, including Apache Commons CLI, Commons CSV, Commons Lang, Google Gson, and additional benchmark projects. For each benchmark, the release includes project manifests and target-method metadata used to define the unit-test generation tasks.&lt;/p&gt;&lt;p dir="ltr"&gt;The experimental-results package contains one directory per experiment run. Each run includes: (1) summary-level effectiveness metrics; (2) per-class and per-target execution records; (3) baseline results without skill injection; (4) results obtained with skill selection and iterative skill evolution; (5) coverage and, where enabled, mutation-testing measurements; (6) execution status and logs; and (7) skill-evolution records. The reported metrics include compilation success, test pass rate, line coverage, branch coverage, mutation score, and token usage.&lt;/p&gt;&lt;p dir="ltr"&gt;The framework compares a no-skill baseline with a skill-enhanced setting. In the enhanced setting, the system selects test-generation skills based on observed failure signals and may evolve skills through iterative generation, execution, feedback, and repair cycles.&lt;/p&gt;&lt;p dir="ltr"&gt;&lt;br&gt;&lt;/p&gt;&lt;p dir="ltr"&gt;The archive is intended to support reproduction, validation, and secondary analysis of experiments on skill-guided LLM-based Java unit-test generation. Generated build artifacts (for example, Maven target directories), credentials, and local configuration overrides are excluded.&lt;/p&gt;</dc:description>
          <dc:date>2026-09-23T11:52:46Z</dc:date>
          <dc:type>Dataset</dc:type>
          <dc:type>Dataset</dc:type>
          <dc:identifier>10.6084/m9.figshare.33972877.v1</dc:identifier>
          <dc:relation>https://figshare.com/articles/dataset/Java_dataset_and_results/33972877</dc:relation>
          <dc:rights>CC BY 4.0</dc:rights>
        </oai_dc:dc>
      </metadata>
    </record>
  </GetRecord>
</OAI-PMH>
