<?xml version='1.0' encoding='utf-8'?>
<?xml-stylesheet type="text/xsl" href="/v2/static/oai2.xsl"?>
<OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd">
  <responseDate>2026-10-07T07:36:02Z</responseDate>
  <request identifier="oai:figshare.com:article/34018209" metadataPrefix="oai_dc" verb="GetRecord">https://api.figshare.com/v2/oai</request>
  <GetRecord>
    <record>
      <header>
        <identifier>oai:figshare.com:article/34018209</identifier>
        <datestamp>2026-09-28T18:56:40Z</datestamp>
        <setSpec>category_29182</setSpec>
        <setSpec>item_type_7</setSpec>
        <setSpec>month_year_09_2026</setSpec>
      </header>
      <metadata>
        <oai_dc:dc xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance"  xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
          <dc:title>An Empirical Investigation of Pre-Trained Deep Learning Model Reuse in the Scientific Process</dc:title>
          <dc:creator>Nicholas Synovic (11655277)</dc:creator>
          <dc:creator>Karolina Ryzka (20591594)</dc:creator>
          <dc:creator>Alessandra Vellucci Solari (20591690)</dc:creator>
          <dc:creator>Kenny Lyons (23549448)</dc:creator>
          <dc:creator>James C. Davis (14565206)</dc:creator>
          <dc:creator>George K. Thiruvathukal (470175)</dc:creator>
          <dc:subject>Empirical software engineering</dc:subject>
          <dc:subject>Pre-trained Models</dc:subject>
          <dc:subject>Software Reuse</dc:subject>
          <dc:subject>Empirical Software Engineering</dc:subject>
          <dc:subject>AI for Science and Engineering</dc:subject>
          <dc:subject>Scientific Process</dc:subject>
          <dc:subject>Large Language Models</dc:subject>
          <dc:subject>Automated Literature Review</dc:subject>
          <dc:subject>Research Productivity</dc:subject>
          <dc:description>&lt;p dir="ltr"&gt;Deep learning has achieved recognition for its impact within natural sciences, yet the prohibitive financial and technical cost of training models from scratch inhibits adoption. Following software engineering community guidance, natural scientists are reusing pre-trained deep learning models (PTMs) to amortize these costs. While prior works recommend PTM reuse patterns, we present the first empirical study of these patterns in the natural sciences, quantifying the utilization and impact of PTM reuse within the scientific process across 17,718 peer reviewed, open access papers. Our results show that “Biochemistry, Genetics and Molecular Biology” has outpaced other natural scientific fields in PTM reuse, “adaptation” reuse is the most prevalent PTM reuse pattern identified across all natural science fields, and the “testing” stage of the scientific process has been most impacted by PTM integration.&lt;/p&gt;</dc:description>
          <dc:date>2026-09-28T18:56:40Z</dc:date>
          <dc:type>Text</dc:type>
          <dc:type>Presentation</dc:type>
          <dc:identifier>10.6084/m9.figshare.34018209.v1</dc:identifier>
          <dc:relation>https://figshare.com/articles/presentation/An_Empirical_Investigation_of_Pre-Trained_Deep_Learning_Model_Reuse_in_the_Scientific_Process/34018209</dc:relation>
          <dc:rights>CC BY 4.0</dc:rights>
        </oai_dc:dc>
      </metadata>
    </record>
  </GetRecord>
</OAI-PMH>
