<?xml version="1.0" encoding="UTF-8"?>
<!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "http://dtd.nlm.nih.gov/publishing/2.0/journalpublishing.dtd">
<article xmlns:xlink="http://www.w3.org/1999/xlink" article-type="research-article" dtd-version="2.0">
  <front>
    <journal-meta>
      <journal-id journal-id-type="publisher-id">JPH</journal-id>
      <journal-id journal-id-type="nlm-ta">JMIR Public Health Surveill</journal-id>
      <journal-title>JMIR Public Health and Surveillance</journal-title>
      <issn pub-type="epub">2369-2960</issn>
      <publisher>
        <publisher-name>JMIR Publications</publisher-name>
        <publisher-loc>Toronto, Canada</publisher-loc>
      </publisher>
    </journal-meta>
    <article-meta>
      <article-id pub-id-type="publisher-id">v12i1e99643</article-id>
      <article-id pub-id-type="pmid">42748426</article-id>
      <article-id pub-id-type="doi">10.2196/99643</article-id>
      <article-categories>
        <subj-group subj-group-type="heading">
          <subject>Original Paper</subject>
        </subj-group>
        <subj-group subj-group-type="article-type">
          <subject>Original Paper</subject>
        </subj-group>
      </article-categories>
      <title-group>
        <article-title>Using Natural Language Processing to Examine State Child Maltreatment Policies and Associations With Outcomes: Multistate Cross-Sectional Study</article-title>
      </title-group>
      <contrib-group>
        <contrib contrib-type="editor">
          <name>
            <surname>Mavragani</surname>
            <given-names>Amaryllis</given-names>
          </name>
        </contrib>
        <contrib contrib-type="editor">
          <name>
            <surname>Sanchez</surname>
            <given-names>Travis</given-names>
          </name>
        </contrib>
      </contrib-group>
      <contrib-group>
        <contrib contrib-type="reviewer">
          <name>
            <surname>Thibodeau</surname>
            <given-names>Eric L</given-names>
          </name>
        </contrib>
        <contrib contrib-type="reviewer">
          <name>
            <surname>Lloyd Sieger</surname>
            <given-names>Margaret</given-names>
          </name>
        </contrib>
        <contrib contrib-type="reviewer">
          <name>
            <surname>Tutsoy</surname>
            <given-names>Onder</given-names>
          </name>
        </contrib>
      </contrib-group>
      <contrib-group>
        <contrib id="contrib1" contrib-type="author" corresp="yes">
          <name name-style="western">
            <surname>Luo</surname>
            <given-names>Zhidi</given-names>
          </name>
          <degrees>PhD</degrees>
          <xref rid="aff1" ref-type="aff">1</xref>
          <address>
            <institution>Department of Psychiatry and Behavioral Sciences</institution>
            <institution>Northwestern University Feinberg School of Medicine</institution>
            <addr-line>Abbott Hall, 12th Floor, 710 N. Lake Shore Dr</addr-line>
            <addr-line>Chicago, IL, 60611</addr-line>
            <country>United States</country>
            <phone>1 312 926 2323</phone>
            <email>zhidi.luo@northwestern.edu</email>
          </address>
          <xref rid="aff2" ref-type="aff">2</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0009-0006-7475-179X</ext-link>
        </contrib>
        <contrib id="contrib2" contrib-type="author">
          <name name-style="western">
            <surname>Epstein</surname>
            <given-names>Richard A</given-names>
          </name>
          <degrees>MPH, PhD</degrees>
          <xref rid="aff1" ref-type="aff">1</xref>
          <xref rid="aff2" ref-type="aff">2</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0001-6624-7311</ext-link>
        </contrib>
        <contrib id="contrib3" contrib-type="author">
          <name name-style="western">
            <surname>Sambamoorthi</surname>
            <given-names>Nethra</given-names>
          </name>
          <degrees>PhD</degrees>
          <xref rid="aff3" ref-type="aff">3</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0002-9949-7306</ext-link>
        </contrib>
        <contrib id="contrib4" contrib-type="author">
          <name name-style="western">
            <surname>Muhammad</surname>
            <given-names>Lutfiyya N</given-names>
          </name>
          <degrees>PhD, MPH</degrees>
          <xref rid="aff4" ref-type="aff">4</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0001-8655-1484</ext-link>
        </contrib>
        <contrib id="contrib5" contrib-type="author">
          <name name-style="western">
            <surname>Jordan</surname>
            <given-names>Neil</given-names>
          </name>
          <degrees>PhD</degrees>
          <xref rid="aff1" ref-type="aff">1</xref>
          <xref rid="aff4" ref-type="aff">4</xref>
          <xref rid="aff5" ref-type="aff">5</xref>
          <ext-link ext-link-type="orcid">https://orcid.org/0000-0001-8467-822X</ext-link>
        </contrib>
      </contrib-group>
      <aff id="aff1">
        <label>1</label>
        <institution>Department of Psychiatry and Behavioral Sciences</institution>
        <institution>Northwestern University Feinberg School of Medicine</institution>
        <addr-line>Chicago, IL</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff2">
        <label>2</label>
        <institution>Health Sciences Integrated Program</institution>
        <institution>Northwestern University Feinberg School of Medicine</institution>
        <addr-line>Chicago, IL</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff3">
        <label>3</label>
        <institution>School of Professional Studies</institution>
        <institution>Northwestern University</institution>
        <addr-line>Chicago, IL</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff4">
        <label>4</label>
        <institution>Department of Preventive Medicine</institution>
        <institution>Northwestern University Feinberg School of Medicine</institution>
        <addr-line>Chicago, IL</addr-line>
        <country>United States</country>
      </aff>
      <aff id="aff5">
        <label>5</label>
        <institution>Department of Medical Social Sciences</institution>
        <institution>Northwestern University Feinberg School of Medicine</institution>
        <addr-line>Chicago, IL</addr-line>
        <country>United States</country>
      </aff>
      <author-notes>
        <corresp>Corresponding Author: Zhidi Luo <email>zhidi.luo@northwestern.edu</email></corresp>
      </author-notes>
      <pub-date pub-type="collection">
        <year>2026</year>
      </pub-date>
      <pub-date pub-type="epub">
        <day>16</day>
        <month>9</month>
        <year>2026</year>
      </pub-date>
      <volume>12</volume>
      <elocation-id>e99643</elocation-id>
      <history>
        <date date-type="received">
          <day>27</day>
          <month>4</month>
          <year>2026</year>
        </date>
        <date date-type="rev-request">
          <day>20</day>
          <month>7</month>
          <year>2026</year>
        </date>
        <date date-type="rev-recd">
          <day>3</day>
          <month>8</month>
          <year>2026</year>
        </date>
        <date date-type="accepted">
          <day>8</day>
          <month>8</month>
          <year>2026</year>
        </date>
      </history>
      <copyright-statement>©Zhidi Luo, Richard A Epstein, Nethra Sambamoorthi, Lutfiyya N Muhammad, Neil Jordan. Originally published in JMIR Public Health and Surveillance (https://publichealth.jmir.org), 16.09.2026.</copyright-statement>
      <copyright-year>2026</copyright-year>
      <license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/">
        <p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (https://creativecommons.org/licenses/by/4.0/), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Public Health and Surveillance, is properly cited. The complete bibliographic information, a link to the original publication on https://publichealth.jmir.org, as well as this copyright and license information must be included.</p>
      </license>
      <self-uri xlink:href="https://publichealth.jmir.org/2026/1/e99643" xlink:type="simple"/>
      <abstract>
        <sec sec-type="background">
          <title>Background</title>
          <p>Child maltreatment is a major public health issue in the United States, with substantial variation in how states define, report, and respond to abuse and neglect. While prior research has examined individual policy components, less is known about how multiple policies jointly shape broader policy environments and relate to maltreatment outcomes.</p>
        </sec>
        <sec sec-type="objective">
          <title>Objective</title>
          <p>This study used natural language processing (NLP) to characterize state child maltreatment policy environments and examine their associations with maltreatment incidence, recurrence, and fatalities across US jurisdictions.</p>
        </sec>
        <sec sec-type="methods">
          <title>Methods</title>
          <p>A cross-sectional study was conducted using 2021 data from 50 US states, the District of Columbia, and Puerto Rico (N=52). A total of 411 state maltreatment policy items were derived from the State Child Abuse &#38; Neglect Policies Database. A total of 6 NLP models (bidirectional and auto-regressive transformer [BART], bidirectional encoder representations from transformers [BERT], robustly optimized BERT approach [RoBERTa], decoding-enhanced BERT with disentangled attention [DeBERTa], Copilot, and LLaMA 3.1) were applied using a zero-shot classification framework to quantify policy characteristics. Models were evaluated using intrinsic (category consistency and semantic alignment) and extrinsic (factor analysis and clustering performance) metrics. Exploratory factor analysis was used to identify latent policy domains, and k-means clustering was applied to group jurisdictions with similar policy profiles. Maltreatment outcomes, including incidence, recurrence, and fatalities, were obtained from the National Child Abuse and Neglect Data System. Outcome differences across clusters were assessed using ANOVA with post hoc pairwise comparisons.</p>
        </sec>
        <sec sec-type="results">
          <title>Results</title>
          <p>The study population included 72,838,819 children and 3,774,528 maltreatment reports, of which 751,283 were substantiated or indicated cases. Across 6 NLP models, DeBERTa demonstrated the best overall performance. Exploratory factor analysis identified 3 primary policy domains: maltreatment definition, mandated reporting, and alternative response. A total of 4 policy clusters were identified. Jurisdictions with weaker reporting requirements and fewer penalties (cluster 2) had the highest incidence of maltreatment (mean 17.51, SD 6.13 vs mean 9.65, SD 4.10; mean 10.28, SD 6.47; and mean 11.74, SD 6.47 per 1000 children; <italic>P</italic>=.07) and recurrence (mean 1.74, SD 0.87 vs mean 0.52, SD 0.44; mean 0.66, SD 0.52; and mean 0.68, SD 0.43 per 1000 children; <italic>P</italic>&#60;.001). Jurisdictions with stronger reporting requirements, broader definitions, and greater use of alternative response systems (cluster 1) had the lowest incidence and recurrence. No statistically significant differences in maltreatment fatalities were observed across clusters (<italic>P</italic>=.89). However, in a subanalysis of 18 fatality-related policy items, clusters differed significantly in fatality rates, with jurisdictions characterized by less clearly defined fatality policies exhibiting higher fatality rates (mean 47.91, SD 8.35 vs mean 22.64, SD 11.82; mean 25.63, SD 18.62; and mean 19.03, SD 11.64 per 1,000,000 children; <italic>P</italic>=.03).</p>
        </sec>
        <sec sec-type="conclusions">
          <title>Conclusions</title>
          <p>State child maltreatment policies form distinct, multidimensional patterns associated with maltreatment outcomes. The findings highlight the importance of evaluating integrated policy frameworks rather than isolated components and demonstrate the use of NLP for large-scale, data-driven policy analysis.</p>
        </sec>
      </abstract>
      <kwd-group>
        <kwd>child abuse</kwd>
        <kwd>child neglect</kwd>
        <kwd>cluster analysis</kwd>
        <kwd>clustering analysis</kwd>
        <kwd>data mining</kwd>
        <kwd>factor analysis</kwd>
        <kwd>health policy</kwd>
        <kwd>natural language processing</kwd>
        <kwd>public policy</kwd>
      </kwd-group>
    </article-meta>
  </front>
  <body>
    <sec sec-type="introduction">
      <title>Introduction</title>
      <p>Child maltreatment, which includes both child abuse and neglect, is a critical public health concern in the United States. In fiscal year 2023, the most recent year for which information is available, there were 546,159 reported children who experienced child maltreatment, with an estimated 2000 children dying from maltreatment [<xref ref-type="bibr" rid="ref1">1</xref>]. Child maltreatment is associated with an increased risk of adverse health outcomes and risk behaviors, including substance use, mental health disorders, chronic physical conditions, and premature mortality [<xref ref-type="bibr" rid="ref2">2</xref>-<xref ref-type="bibr" rid="ref5">5</xref>]. The economic burden is substantial, with the total lifetime cost estimated at US $2.94 trillion for investigated cases and US $563 billion for substantiated cases in the federal fiscal year 2018 [<xref ref-type="bibr" rid="ref6">6</xref>].</p>
      <p>Child maltreatment policy refers to states’ definitions of child abuse and neglect, along with related regulations that shape child welfare practices. These policies determine how allegations are referred, reported, investigated, and managed [<xref ref-type="bibr" rid="ref7">7</xref>]. For example, mandated reporting expansion refers to policies that broaden the categories of individuals required to report suspected maltreatment (eg, extending beyond professionals such as teachers and health care providers to include additional occupational groups or, in some states, all adults) [<xref ref-type="bibr" rid="ref8">8</xref>]. Similarly, policies such as centralized intake systems standardize how reports are received and screened, while differential response systems allow for alternative, noninvestigative pathways for lower-risk cases [<xref ref-type="bibr" rid="ref9">9</xref>]. Together, these policy components influence both the volume and characteristics of cases that enter and progress through the child welfare system.</p>
      <p>At the federal level, the Child Abuse Prevention and Treatment Act (CAPTA), first enacted in 1974 and reauthorized multiple times since, establishes minimum standards for state definitions of abuse and neglect and conditions federal funding on states’ compliance with these requirements [<xref ref-type="bibr" rid="ref10">10</xref>]. However, states retain substantial discretion in translating these federal requirements into policy and practice, resulting in significant cross-state variation [<xref ref-type="bibr" rid="ref7">7</xref>]. This variation leads to substantial differences in reported maltreatment rates and intervention effectiveness, with substantiation rates ranging from 5% to 45% across states [<xref ref-type="bibr" rid="ref11">11</xref>].</p>
      <p>Research suggests that states spending an additional US $1000 per person in poverty on benefits such as cash aid, child care, or Medicaid assistance saw 4.3% fewer maltreatment reports and 4.0% fewer substantiations [<xref ref-type="bibr" rid="ref12">12</xref>]. Additionally, states that implemented mandated reporting expansion, centralized intake, and increased state staffing saw a 32% increase in maltreatment reports, although substantiated reports declined by 5% to 6% [<xref ref-type="bibr" rid="ref7">7</xref>]. In contrast, states adopting both differential response and higher standards of proof experienced a 24% decrease in substantiated reports, with differential response accounting for 11% and higher proof standards contributing between 12% and 13% [<xref ref-type="bibr" rid="ref7">7</xref>]. A separate national quasi-experimental study using data from 2004 to 2017 found that states with differential response programs had approximately 19% fewer substantiated reports and a 17% reduction in foster care use relative to states without differential response [<xref ref-type="bibr" rid="ref9">9</xref>]. These literatures reflect the complex and sometimes countervailing ways in which policy components interact with one another and with local implementation contexts.</p>
      <p>However, existing studies have largely examined individual policies or small sets of policy features in isolation. Less is known about how multiple policy components co-occur within states to form broader policy environments, or how states differ systematically across key policy domains. This limitation is important because child welfare systems operate as integrated structures, where policies governing reporting, screening, investigation, and service response interact with one another. As a result, the relationship between any single policy component and observed outcomes may depend on the broader policy context in which it is embedded.</p>
      <p>Recent advancements in natural language processing (NLP), particularly the development of large language models (LLMs), have begun to transform policy research and child welfare research. In policy text mining, zero-shot and few-shot classification have been applied to large corpora of legislative text, demonstrating scalability without extensive hand-labeled training data [<xref ref-type="bibr" rid="ref13">13</xref>]. Other work has used NLP to construct interpretable, theory-driven measures of legislative quality from statutory text [<xref ref-type="bibr" rid="ref14">14</xref>] and to track policy status across large sets of climate and environmental documents [<xref ref-type="bibr" rid="ref15">15</xref>-<xref ref-type="bibr" rid="ref17">17</xref>]. In child welfare research, NLP and related methods have focused on case-level rather than policy-level data. Predictive risk modeling using structured administrative records has estimated individual maltreatment risk, in some cases outperforming existing screening tools [<xref ref-type="bibr" rid="ref18">18</xref>,<xref ref-type="bibr" rid="ref19">19</xref>], while NLP applied to case notes and health records has identified instances of maltreatment [<xref ref-type="bibr" rid="ref20">20</xref>-<xref ref-type="bibr" rid="ref23">23</xref>]. These efforts enhanced data-driven decision-making within child welfare agencies by allowing for more nuanced insights into family dynamics and case histories [<xref ref-type="bibr" rid="ref24">24</xref>]. These approaches demonstrate NLP’s value for surfacing information from unstructured text, but case note narratives have also been shown to embed caseworker bias and uncertainty, raising caution about their use in downstream predictive tasks [<xref ref-type="bibr" rid="ref25">25</xref>]. However, none of this literature has applied NLP to the formal statutory and regulatory policy documents that define state child welfare systems, as distinct from case-level clinical or casework text.</p>
      <p>This study uses NLP to identify key components of state child maltreatment policies, examine how states differ across these components, and assess how these policy configurations are associated with maltreatment outcomes in the United States. By characterizing the multidimensional structure of state policy environments, this study aims to provide a more comprehensive understanding of how policy variation corresponds to differences in observed child welfare outcomes.</p>
    </sec>
    <sec sec-type="methods">
      <title>Methods</title>
      <sec>
        <title>Overview</title>
        <p>This study was reported in accordance with the STROBE (Strengthening the Reporting of Observational Studies in Epidemiology) guidelines. The overall analysis proceeded through 3 stages, as illustrated in <xref rid="figure1" ref-type="fig">Figure 1</xref>: comparing and evaluating 6 NLP models and quantifying the 411 state policy items (“Quantification of State Maltreatment Policy Using NLP” section), identifying latent policy dimensions via factor analysis and jurisdiction grouping via k-means clustering (“Exploratory Factor Analysis” and “Clustering Analysis” sections), and comparison of demographic, economic, and maltreatment outcomes across the resulting policy clusters (“Demographic, Economic, and Outcome Comparisons” section).</p>
        <fig id="figure1" position="float">
          <label>Figure 1</label>
          <caption>
            <p>Analytical workflow for evaluating state child maltreatment policy configurations using natural language processing (NLP), exploratory factor analysis, and clustering among 52 US jurisdictions in 2021. ACS: American Community Survey; BART: bidirectional and auto-regressive transformer; BERT: bidirectional encoder representations from transformers; BIC: Bayesian information criterion; BSS/WSS: between-cluster sum of squares to within-cluster sum of squares ratio; DeBERTa: decoding-enhanced BERT with disentangled attention; MRSF: multiple R-squared of scores with factors; NCANDS: National Child Abuse and Neglect Data System; RoBERTa: robustly optimized BERT approach.</p>
          </caption>
          <graphic xlink:href="publichealth_v12i1e99643_fig1.png" alt-version="no" mimetype="image" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec>
        <title>Data Sources</title>
        <p>This study used data from 3 sources. State-level child abuse and neglect policy data from 2021 were obtained from the State Child Abuse &#38; Neglect Policies Database (SCAN), which documents state child welfare laws and policies across all 50 US states, the District of Columbia, and Puerto Rico. The database contains 411 standardized policy variables covering state-specific definitions of child abuse and neglect, reporting policies, screening policies, investigation policies, child welfare responses, child welfare system context, and child fatality–related policies in calendar year 2021 [<xref ref-type="bibr" rid="ref26">26</xref>]. SCAN captures written state laws and policies from publicly available or state-provided sources, including state statutes, regulations, and agency policy documents. Policy variables are coded using a standardized protocol and are verified by state agency contacts when feasible.</p>
        <p>Additionally, child welfare agency– and child-level maltreatment data from 2021 were sourced from the National Child Abuse and Neglect Data System (NCANDS), a federally sponsored program administered by the US Department of Health and Human Services [<xref ref-type="bibr" rid="ref27">27</xref>,<xref ref-type="bibr" rid="ref28">28</xref>]. NCANDS collects and analyzes state-reported data on child abuse and neglect known to child protective services, as mandated by the CAPTA [<xref ref-type="bibr" rid="ref27">27</xref>,<xref ref-type="bibr" rid="ref28">28</xref>]. Child-level records were used to derive state-level measures of maltreatment incidence and recurrence by aggregating individual cases. Maltreatment fatality rates were obtained from the NCANDS agency-level data because some fatality information is not available in the child-level files for confidentiality and data security reasons.</p>
        <p>Finally, demographic and economic data from 2021 were drawn from the American Community Survey, an annual survey conducted by the US Census Bureau that provides detailed estimates of social, economic, housing, and demographic characteristics at the state and local levels [<xref ref-type="bibr" rid="ref29">29</xref>].</p>
      </sec>
      <sec>
        <title>Study Cohort and Outcomes</title>
        <p>The study cohort included all 50 states in the United States, the District of Columbia, and the Commonwealth of Puerto Rico, with 411 state maltreatment policy items assessed in each jurisdiction in calendar year 2021. Each policy item represents a specific child welfare policy provision (eg, a mandated reporting requirement, definition of child maltreatment, screening procedure, or investigation practice) and is organized within the SCAN database into predefined policy domains, including child maltreatment definitions, mandated reporting, screening, investigations, child welfare responses, and child welfare system context (Table S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). These domains are part of the original SCAN database structure. Policy items vary in measurement format depending on the underlying policy provision and include binary, categorical, ordinal, and descriptive variables. Although the SCAN database organizes policy items into predefined domains, these domains were used only to describe the source variables. The subsequent analyses identified latent policy dimensions empirically using exploratory factor analysis (EFA) rather than relying on the original SCAN categorization.</p>
        <p>Maltreatment outcomes were evaluated for individuals aged &#60;18 years. While some states, such as Illinois, allow youth aged 18 years or older to remain in care, others do not. To ensure consistent denominators across jurisdictions, individuals aged 18 years or older were excluded.</p>
        <p>The study examined three outcomes: (1) incidence of maltreatment, measured as the number of substantiated or indicated maltreatment cases per 1000 children aged &#60;18 years in each state; (2) maltreatment recurrence, defined as the number of children previously identified as having experienced maltreatment who experienced subsequent incidents per 1000 children in the population; and (3) maltreatment fatalities, defined as the number of deaths resulting from maltreatment per 1,000,000 children. For the first 2 outcomes, only substantiated or indicated maltreatment recorded by state child protective services in NCANDS with report dates in 2021 were included. Although these outcomes represent related dimensions of child maltreatment burden, they capture distinct aspects of child welfare involvement and were analyzed separately rather than modeled as a combined outcome.</p>
      </sec>
      <sec>
        <title>Quantification of State Maltreatment Policy Using NLP</title>
        <sec>
          <title>Overview</title>
          <p>The primary study variable was state maltreatment policy. To prepare these data, 411 state maltreatment policy items from 2021 across 52 US jurisdictions were processed using NLP models.</p>
        </sec>
        <sec>
          <title>NLP Approach</title>
          <p>This study applied NLP methods to classify and analyze 411 state maltreatment policy items from 2021 across 52 US jurisdictions. A zero-shot classification [<xref ref-type="bibr" rid="ref30">30</xref>], which uses pretrained models to assign labels without task-specific fine-tuning, was selected because the policy data were not labeled for the specific classification tasks required in this study. Alternative approaches, such as supervised or few-shot classification, typically require curated training examples for each category, which were not available and would require substantial manual annotation. Zero-shot methods provide a scalable and consistent framework for applying predefined policy categories across heterogeneous statutory text while avoiding potential bias introduced through hand-labeled training data.</p>
          <p>To evaluate the robustness of this approach, 6 NLP models were compared, representing distinct architectures and training paradigms, including encoder-based, sequence-to-sequence, and generative LLMs. Comparing these models allows assessment of whether classification results are consistent across fundamentally different modeling approaches, and whether more recent generative models provide advantages over traditional transformer-based classifiers in capturing the structure of policy text. The 6 models evaluated were as follows:</p>
          <list list-type="order">
            <list-item>
              <p>Bidirectional and auto-regressive transformer (BART) [<xref ref-type="bibr" rid="ref31">31</xref>], a sequence-to-sequence transformer with a bidirectional encoder and an autoregressive decoder, designed for classification and text generation tasks;</p>
            </list-item>
            <list-item>
              <p>Bidirectional encoder representations from transformers (BERT) [<xref ref-type="bibr" rid="ref32">32</xref>], a bidirectional transformer encoder trained using masked-language modeling for deep contextual text understanding;</p>
            </list-item>
            <list-item>
              <p>Robustly optimized BERT approach (RoBERTa) [<xref ref-type="bibr" rid="ref33">33</xref>], an optimized version of BERT trained longer and on larger corpora to improve robustness and performance;</p>
            </list-item>
            <list-item>
              <p>Decoding-enhanced BERT with disentangled attention (DeBERTa) [<xref ref-type="bibr" rid="ref34">34</xref>], a transformer model that enhances BERT/RoBERTa with disentangled attention and improved position encoding;</p>
            </list-item>
            <list-item>
              <p>Microsoft Copilot [<xref ref-type="bibr" rid="ref35">35</xref>], generative transformer models adapted for instruction-following and classification;</p>
            </list-item>
            <list-item>
              <p>Meta LLaMA 3.1 [<xref ref-type="bibr" rid="ref36">36</xref>], an LLM in the LLaMA family, trained as an autoregressive transformer with strong general-purpose language understanding capabilities.</p>
            </list-item>
          </list>
          <p>The models evaluated each policy item independently by assessing whether each jurisdiction’s policy text supports or contradicts the policy item’s title description. For example, for the policy item titled “How is reporting decentralized?,” jurisdiction responses include “Each county has its own reporting hotline,” “Some counties have their own reporting hotlines,” and “No county has its own reporting hotline.” Each NLP model assigned a likelihood score from 0 to 1 to each response, indicating the degree to which it supported the policy title, where 1 represented the strongest support and 0 the weakest. For instance, the BART model assigned scores of 0.90, 0.81, and 0.59, respectively.</p>
          <p>For LLMs, Microsoft Copilot and Meta LLaMA 3.1, the scores were generated by a structured prompt, which was applied separately to each policy item and took the following format:</p>
          <p>“Please estimate the likelihood (on a scale from 0 to 1) that the response [A], [B], [C] supports or opposes the statement [Policy Title]. Stronger support should be closer to 1, a neutral or irrelevant stance should be around 0.5, and complete opposition should be close to 0. Provide likelihoods in this format: [A]: score1; [B]: score2; [C]: score3.”</p>
        </sec>
        <sec>
          <title>NLP Evaluation Metrics</title>
          <p>The 6 models were evaluated using intrinsic and extrinsic metrics [<xref ref-type="bibr" rid="ref37">37</xref>,<xref ref-type="bibr" rid="ref38">38</xref>]. Intrinsic evaluation measured internal consistency and agreement by testing how well LLM embedding vectors capture the underlying characteristics of the texts, while extrinsic evaluation assessed the performance of key embedding features extracted from the policy documents based on their performance in downstream tasks and analyses [<xref ref-type="bibr" rid="ref37">37</xref>]. Since later sections of the study applied EFA [<xref ref-type="bibr" rid="ref39">39</xref>] and k-means clustering [<xref ref-type="bibr" rid="ref40">40</xref>] to policy data, these techniques were used here as extrinsic evaluation metrics to determine how well different NLP models capture meaningful variation in policy content for the same set of features.</p>
          <p>For intrinsic evaluation, 3 metrics were used. First, category consistency was assessed based on the relative ordering of classification scores across response categories. Specifically, items labeled as “Yes” were expected to receive higher scores than those labeled “No,” while intermediate responses such as “Not Applicable,” “Unknown,” or “Missing” were expected to fall between “Yes” and “No.” Second, for descriptive responses, semantic alignment was measured using 2 metrics. Because a true gold standard was unavailable for these responses, 2 proxy gold standards were used: (1) cosine similarity scores, where higher values indicate better semantic alignment, and (2) the average score across all methods. The correlation between each model’s results and these proxy standards, where a higher correlation indicates better agreement, served as the metric to assess semantic alignment.</p>
          <p>Extrinsic evaluation assessed EFA and clustering results derived from each NLP model-specific policy representation. In this stage, each model produced a complete set of scores for the 411 policy items across jurisdictions, which served as inputs for EFA and k-means clustering. Model performance was then assessed by comparing goodness-of-fit and clustering metrics to identify which NLP approaches most effectively capture meaningful structure in the policy data. EFA results were evaluated using the Bayesian information criterion (BIC) [<xref ref-type="bibr" rid="ref41">41</xref>], where lower values indicate better model fit accounting for complexity, and the multiple R-squared of scores with factors (MRSF) [<xref ref-type="bibr" rid="ref39">39</xref>], where higher values reflect better explanatory power of the factor scores. To explore different latent structures, EFA models with 2, 3, and 4 factors were examined. Clustering performance was analyzed using the Calinski-Harabasz score (CH) [<xref ref-type="bibr" rid="ref42">42</xref>], where higher values indicate better cluster separation relative to complexity, and the between-cluster sum of squares to within-cluster sum of squares (BSS/WSS) ratio [<xref ref-type="bibr" rid="ref43">43</xref>], where higher ratios represent more compact and well-separated clusters. Models with 2, 3, and 4 clusters were examined to evaluate clustering performance.</p>
        </sec>
        <sec>
          <title>Qualitative Review</title>
          <p>Finally, a qualitative review explored why certain methods perform better on specific metrics by analyzing model outputs, identifying recurring error patterns, and assessing consistency across similar policy items. This review considered factors such as sensitivity to nuanced language, handling of ambiguous or incomplete responses, and stability in scoring across different policy categories. This review aimed to gain an intuitive understanding of model behavior and to verify that the results aligned with expectations.</p>
          <p>Overall model performance was determined by considering results across all intrinsic and extrinsic evaluation metrics rather than relying on a single metric or composite score. No weighting scheme was applied. Instead, models were compared based on their relative performance patterns across evaluation dimensions, with qualitative review used to further examine model behavior, identify potential sources of variation, and assess whether quantitative findings were consistent with expected policy interpretation.</p>
        </sec>
      </sec>
      <sec>
        <title>Covariates</title>
        <p>Demographic and economic characteristics were included as covariates to contextualize jurisdiction-level differences. Demographic covariates included age distribution, racial and ethnic composition, and gender. Economic conditions were measured by the proportion of the population living below 50%, 100%, 125%, and 400% of the federal poverty level, capturing varying levels of economic vulnerability.</p>
      </sec>
      <sec>
        <title>Statistical Analysis</title>
        <sec>
          <title>Exploratory Factor Analysis</title>
          <p>Although the SCAN database organizes policy items into predefined domains for documentation purposes, these domains do not necessarily represent latent policy constructs. EFA was conducted on the 411 state maltreatment policy items to identify latent dimensions [<xref ref-type="bibr" rid="ref39">39</xref>]. Consequently, policy items originating from the same SCAN domain were not constrained to load on the same factor; likewise, policy items from different SCAN domains could load on the same factor if they reflected similar underlying policy characteristics. The number of factors was determined using a scree plot, a graphical tool that helps identify the point where the rate of decrease in explained variance slows [<xref ref-type="bibr" rid="ref44">44</xref>]. While this plot provides a helpful visual cue, identifying the “elbow” can be subjective and challenging [<xref ref-type="bibr" rid="ref44">44</xref>,<xref ref-type="bibr" rid="ref45">45</xref>]. In cases where the exact location of the “elbow” is ambiguous and several factor number choices were similarly justifiable, we opted for the smaller number of factors to ensure concise and easier interpretation.</p>
          <p>To identify the underlying themes for each factor, policy items with high factor loadings were reviewed. These factors were then labeled based on the substantive patterns they capture. For each factor in each jurisdiction, factor scores were computed to quantify the degree to which each jurisdiction aligns with the identified policy dimensions [<xref ref-type="bibr" rid="ref46">46</xref>]. The factor scores provided a reduced representation of the original 411 policy items and were used in subsequent analyses to characterize jurisdiction-level policy configurations.</p>
          <p>The internal consistency of policy items contributing to each identified factor was assessed using McDonald ω [<xref ref-type="bibr" rid="ref47">47</xref>]. ω was calculated based on policy items with absolute factor loadings above 0.4.</p>
        </sec>
        <sec>
          <title>Clustering Analysis</title>
          <p>K-means clustering was applied to the 411 state maltreatment policy items to identify groups of jurisdictions that share similar policy characteristics [<xref ref-type="bibr" rid="ref40">40</xref>]. The optimal number of clusters was determined using the Silhouette method [<xref ref-type="bibr" rid="ref48">48</xref>], which evaluated cluster cohesion and separation, along with interpretability considerations, to ensure the resulting groupings were meaningful and practically relevant. After the number of clusters is selected, each jurisdiction is assigned to a cluster based on the similarity of its policy item representations to the identified cluster profiles. To visually illustrate the geographic distribution of policy clusters, a US map is generated, with jurisdictions color-coded according to their assigned cluster. The policy characteristics of each cluster are identified by comparing average factor scores derived from the factor analysis.</p>
          <p>Although k-means clustering and factor analysis are distinct techniques, they serve complementary purposes in this study. Factor analysis reduces the dimensionality of the policy items and identifies underlying policy domains, whereas clustering groups jurisdictions based on similarity across the full set of policy features. By conducting clustering on all items and using factor scores for interpretation, this approach preserves the full information available for grouping while enabling a more interpretable, lower-dimensional characterization of cluster differences. Consistency between cluster profiles and factor structures further supports the coherence and robustness of the observed policy patterns.</p>
        </sec>
        <sec>
          <title>Demographic, Economic, and Outcome Comparisons</title>
          <p>To characterize each cluster, policy factor scores were compared across clusters. Then, demographic, economic, and outcome differences were examined across jurisdiction clusters. Statistical differences across clusters were tested using ANOVA for continuous variables [<xref ref-type="bibr" rid="ref49">49</xref>] and chi-square tests for categorical variables [<xref ref-type="bibr" rid="ref50">50</xref>]. When significant differences are detected, post hoc pairwise comparisons are conducted to determine which clusters differ from each other. Given the small sample size of 52 US jurisdictions, a significance level of 0.10 is used to reduce the risk of type II errors and detect potential associations that may warrant further investigation [<xref ref-type="bibr" rid="ref51">51</xref>,<xref ref-type="bibr" rid="ref52">52</xref>].</p>
        </sec>
      </sec>
      <sec>
        <title>Sensitivity Analyses</title>
        <sec>
          <title>Leave-One-State-Out Sensitivity Analysis</title>
          <p>Cluster robustness was assessed using a leave-one-jurisdiction-out sensitivity analysis. K-means clustering was repeated after sequentially removing each jurisdiction, and agreement between the resulting cluster assignments and the original clustering was evaluated using the adjusted Rand index (ARI) [<xref ref-type="bibr" rid="ref53">53</xref>]. This analysis was conducted to evaluate whether the identified policy configurations were robust to the exclusion of individual jurisdictions and not driven by any single state’s policy profile.</p>
        </sec>
        <sec>
          <title>Sensitivity Analysis Including All NCANDS Maltreatment Records</title>
          <p>To assess whether findings were sensitive to the definition of maltreatment outcomes, a sensitivity analysis was conducted using all screened-in maltreatment records reported in NCANDS in 2021, regardless of substantiation or indication status. The same outcome comparison procedures across policy clusters were applied to assess whether the observed associations between policy configurations and maltreatment outcomes remained consistent under an alternative maltreatment definition.</p>
        </sec>
      </sec>
      <sec>
        <title>Subanalysis</title>
        <p>To further examine policies specifically related to maltreatment fatalities, a separate and independent clustering analysis was conducted using only the 18 fatality-related policy items (Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Unlike the primary analysis, which clustered jurisdictions based on all 411 policy items, this subanalysis clustered jurisdictions solely according to fatality-related policies.</p>
        <p>K-means clustering is applied to these 18 policy items to group jurisdictions with similar fatality-related policy characteristics. The maltreatment fatality rates across these clusters are then compared to identify any notable differences related to policy configurations. A violin plot—a visualization that combines a box plot with a density curve—is used to display the distribution of fatality rates by cluster. Statistical tests are conducted to determine whether the observed differences in fatality rates are statistically significant.</p>
      </sec>
      <sec>
        <title>Ethical Considerations</title>
        <p>This study was approved by the Northwestern University Institutional Review Board (STU00222024). The study consisted of a secondary analysis of deidentified administrative data obtained from the National Data Archive on Child Abuse and Neglect (NDACAN). Data access was granted by NDACAN under its Data Use Agreement, and all analyses were conducted in accordance with the terms of the Data Use Agreement and applicable institutional policies. Because the study involved retrospective analysis of existing deidentified data with no direct participant contact, the Northwestern University Institutional Review Board waived the requirement for informed consent.</p>
        <p>The authors did not have access to information that could directly or indirectly identify individual participants during or after data collection. All datasets provided by NDACAN were deidentified prior to release, and analyses were conducted using secure data management procedures in compliance with the NDACAN Data Use Agreement, which prohibits attempts to reidentify participants.</p>
        <p>No participants were recruited for this secondary data analysis; therefore, no compensation was provided.</p>
      </sec>
    </sec>
    <sec sec-type="results">
      <title>Results</title>
      <sec>
        <title>Study Cohort Characteristics</title>
        <p>The study included 52 jurisdictions in the United States: all 50 states, the District of Columbia, and the Commonwealth of Puerto Rico. In calendar year 2021, these jurisdictions had a combined population of 72,838,819 individuals aged &#60;18 years, with 18,416,713 (25.3%) aged &#60;5 years, and 54,422,106 (74.7%) between ages 5 and 17 years. In the same year, there were 3,774,528 maltreatment records, each representing a unique report and child combination, involving 3,134,446 unique children. Of these, 751,283 records were indicated, affecting 692,891 unique children. Although NCANDS contains individual-level child welfare records, the unit of analysis for this study was the jurisdiction. Individual-level records were aggregated to the jurisdiction level to derive maltreatment outcomes.</p>
      </sec>
      <sec>
        <title>NLP Results</title>
        <p><xref ref-type="table" rid="table1">Table 1</xref> presents the comparison of the 6 NLP models: BART, BERT, RoBERTa, DeBERTa, Copilot, and LLaMA 3.1. For intrinsic evaluation, category consistency scores indicate that DeBERTa and Copilot achieve the best performance (0.95). Semantic alignment results vary across metrics. Based on the similarity metric, LLaMA 3.1 performs best (0.21), followed by BART (0.19) and DeBERTa (0.14). In contrast, when evaluated using the average score metric, Copilot ranks highest (0.91).</p>
        <p>For extrinsic evaluation, RoBERTa demonstrates the best performance in factor analysis across all models, achieving the lowest BIC values and the highest MRSF scores for different numbers of factors (2-factor: BIC=–601336; MRSF=0.98; 3-factor: BIC=–561327; MRSF=0.98; 4-factor: BIC=–535315; MRSF=0.98). DeBERTa follows closely with competitive results (2-factor: BIC=–581536; MRSF=0.98; 3-factor: BIC=–550143; MRSF=0.98; 4-factor: BIC=–526241; MRSF=0.97). For clustering, DeBERTa outperforms other models, achieving the highest CH and BSS/WSS ratios across all cluster solutions (2 clusters: CH=4.45; BSS/WSS=1.07; 3 clusters: CH=3.90; BSS/WSS=1.08; 4 clusters: CH=3.28; BSS/WSS=1.09).</p>
        <p>The qualitative review showed key sources of error and model-specific challenges (Table S3 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). For category consistency, RoBERTa had the lowest score due to struggles with negations (eg, “does not” or “failure to report”), leading to misclassifications, while LLaMA 3.1 often hallucinated artificial logic. In semantic alignment, LLM performance varied—while these models sometimes fabricated justifications instead of following response structures, they also demonstrated a better understanding of longer descriptive texts. However, evaluating semantic alignment was challenging, as similarity and average scores served as imperfect proxies. When comparing similarity and average score approaches, DeBERTa was neither the best nor the worst, performing well in some contexts but falling short in others. Finally, in extrinsic evaluation, DeBERTa outperformed other models, likely because it better preserves structural relationships and maintains meaningful variance in the data.</p>
        <p>Based on intrinsic, extrinsic, and qualitative evaluations, DeBERTa was the best-performing model overall, achieving the highest category consistency, strong semantic alignment, the second-best factor analysis performance, and the best clustering result. Because a gold-standard annotation of policy items did not exist, a formal error rate was not calculated.</p>
        <table-wrap position="float" id="table1">
          <label>Table 1</label>
          <caption>
            <p>Comparison of 6 natural language processing methods for quantifying state child maltreatment policies across 52 US jurisdictions using 2021 State Child Abuse &#38; Neglect Policies Database data.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="30"/>
            <col width="210"/>
            <col width="0"/>
            <col width="120"/>
            <col width="0"/>
            <col width="120"/>
            <col width="0"/>
            <col width="140"/>
            <col width="0"/>
            <col width="140"/>
            <col width="0"/>
            <col width="120"/>
            <col width="0"/>
            <col width="120"/>
            <thead>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td colspan="2">BART<sup>a</sup></td>
                <td colspan="2">BERT<sup>b</sup></td>
                <td colspan="2">RoBERTa<sup>c</sup></td>
                <td colspan="2">DeBERTa<sup>d</sup></td>
                <td colspan="2">Copilot</td>
                <td>LLaMA</td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td colspan="14">Intrinsic</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Category consistency<sup>e</sup></td>
                <td colspan="2">0.90</td>
                <td colspan="2">0.81</td>
                <td colspan="2">0.73</td>
                <td colspan="2">0.95</td>
                <td colspan="2">0.95</td>
                <td colspan="2">0.78</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Semantic alignment—similarity<sup>f</sup></td>
                <td colspan="2">0.19</td>
                <td colspan="2">0.04</td>
                <td colspan="2">0.10</td>
                <td colspan="2">0.14</td>
                <td colspan="2">0.10</td>
                <td colspan="2">0.21</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Semantic alignment—average<sup>g</sup></td>
                <td colspan="2">0.77</td>
                <td colspan="2">0.80</td>
                <td colspan="2">0.35</td>
                <td colspan="2">0.68</td>
                <td colspan="2">0.91</td>
                <td colspan="2">0.71</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Extrinsic: factor analysis</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 2, BIC<sup>h,i</sup></td>
                <td colspan="2">–464,493</td>
                <td colspan="2">–460,680</td>
                <td colspan="2">–601,336</td>
                <td colspan="2">–581,536</td>
                <td colspan="2">–457,318</td>
                <td colspan="2">–442,535</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 2, MRSF<sup>i</sup></td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.98</td>
                <td colspan="2">0.98</td>
                <td colspan="2">0.96</td>
                <td colspan="2">0.97</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 3, BIC<sup>i</sup></td>
                <td colspan="2">–439,225</td>
                <td colspan="2">–439,606</td>
                <td colspan="2">–561,327</td>
                <td colspan="2">–550,143</td>
                <td colspan="2">–437,013</td>
                <td colspan="2">–425,035</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 3, MRSF<sup>i</sup></td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.98</td>
                <td colspan="2">0.98</td>
                <td colspan="2">0.96</td>
                <td colspan="2">0.97</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 4, BIC<sup>i</sup></td>
                <td colspan="2">–420,815</td>
                <td colspan="2">–423,379</td>
                <td colspan="2">–535,315</td>
                <td colspan="2">–526,241</td>
                <td colspan="2">–419,011</td>
                <td colspan="2">–409,666</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Factor analysis 4, MRSF<sup>i</sup></td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.98</td>
                <td colspan="2">0.97</td>
                <td colspan="2">0.96</td>
                <td colspan="2">0.96</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Extrinsic: k-means</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 2, CH<sup>i</sup></td>
                <td colspan="2">3.57</td>
                <td colspan="2">2.81</td>
                <td colspan="2">2.67</td>
                <td colspan="2">4.45</td>
                <td colspan="2">2.69</td>
                <td colspan="2">2.81</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 2, BSS/WSS<sup>i</sup></td>
                <td colspan="2">1.06</td>
                <td colspan="2">1.04</td>
                <td colspan="2">1.04</td>
                <td colspan="2">1.07</td>
                <td colspan="2">1.03</td>
                <td colspan="2">1.03</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 3, CH<sup>i</sup></td>
                <td colspan="2">3.07</td>
                <td colspan="2">2.06</td>
                <td colspan="2">2.76</td>
                <td colspan="2">3.90</td>
                <td colspan="2">2.11</td>
                <td colspan="2">2.65</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 3, BSS/WSS<sup>i</sup></td>
                <td colspan="2">1.07</td>
                <td colspan="2">1.03</td>
                <td colspan="2">1.05</td>
                <td colspan="2">1.08</td>
                <td colspan="2">1.03</td>
                <td colspan="2">1.05</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 4, CH<sup>i</sup></td>
                <td colspan="2">2.80</td>
                <td colspan="2">2.49</td>
                <td colspan="2">2.74</td>
                <td colspan="2">3.28</td>
                <td colspan="2">2.04</td>
                <td colspan="2">2.79</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>K-means 4, BSS/WSS<sup>i</sup></td>
                <td colspan="2">1.07</td>
                <td colspan="2">1.06</td>
                <td colspan="2">1.07</td>
                <td colspan="2">1.09</td>
                <td colspan="2">1.04</td>
                <td colspan="2">1.07</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table1fn1">
              <p><sup>a</sup>BART: bidirectional and auto-regressive transformer.</p>
            </fn>
            <fn id="table1fn2">
              <p><sup>b</sup>BERT: bidirectional encoder representations from transformers.</p>
            </fn>
            <fn id="table1fn3">
              <p><sup>c</sup>RoBERTa: robustly optimized BERT approach.</p>
            </fn>
            <fn id="table1fn4">
              <p><sup>d</sup>DeBERTa: decoding-enhanced BERT with disentangled attention.</p>
            </fn>
            <fn id="table1fn5">
              <p><sup>e</sup>Category consistency: evaluates whether categorical descriptions align with expected patterns. Higher values indicate better consistency.</p>
            </fn>
            <fn id="table1fn6">
              <p><sup>f</sup>Semantic alignment—similarity: using a similarity score as the proxy gold standard to evaluate semantic alignment. Higher values indicate better semantic alignment.</p>
            </fn>
            <fn id="table1fn7">
              <p><sup>g</sup>Semantic alignment—average: using the average score across all methods as the proxy gold standard to evaluate semantic alignment. Higher values indicate better semantic alignment.</p>
            </fn>
            <fn id="table1fn8">
              <p><sup>h</sup>BIC: Bayesian information criterion.</p>
            </fn>
            <fn id="table1fn9">
              <p><sup>i</sup>Extrinsic evaluation metric row titles follow the format “[Method] [Number], [Metric],” where the number represents either the number of factors in exploratory factor analysis or the number of clusters in k-means clustering. The metrics include BIC, multiple R-squared of scores with factors (MRSF), Calinski-Harabasz score (CH), and the between-cluster sum of squares to within-cluster sum of squares ratio (BSS/WSS). Lower BIC values indicate a better fit, while higher values are preferable for MRSF, CH, and BSS/WSS.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
      </sec>
      <sec>
        <title>EFA Results</title>
        <p>The scree plot (Figure S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) showed that although a clear “elbow” did not appear until over 50 factors—an impractically large number for interpretation—the decline in eigenvalues slowed around 3 factors, suggesting a point of diminishing returns. We interpreted this inflection as a minor elbow and selected a 3-factor solution to avoid the complexity of higher-dimensional models while still retaining meaningful structure.</p>
        <p>The items with the highest absolute factor loadings for each factor (Table S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) indicated that Factor 1 represented the definition domain, Factor 2 focused on the reporting domain, and Factor 3 captured the alternative response domain. Factor 1 (Definition) reflected the scope of maltreatment definitions, with higher scores corresponding to broader definitions that explicitly included emotional harm, abandonment, and sexual abuse, as well as more stringent safe haven requirements. Factor 2 (Reporting) covered policies related to mandated reporting, where higher scores were associated with a greater number of mandated reporters, more stringent penalties for failure to report, and a reduced emphasis on reporter training. Factor 3 (Alternative Response) captured states’ implementation of differential response systems, with higher scores indicating broader eligibility for alternative responses based on case characteristics such as risk level and maltreatment type.</p>
        <p>The extracted policy dimensions demonstrated high internal consistency. McDonald ω values were 0.996, 0.959, and 0.938 for the 3 factors, indicating strong coherence among policy items contributing to each latent policy dimension.</p>
      </sec>
      <sec>
        <title>Clustering Results</title>
        <p>The k-means clustering analysis was used to identify distinct clusters based on policy characteristics. The optimal number of clusters was determined using the silhouette width, which peaked at 2 and 4 clusters (Figure S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). To better capture the diversity and variability in the policy characteristics across jurisdictions, 4 clusters were chosen.</p>
        <p>A total of 4 distinct policy clusters with varying geographic distributions across US states were identified (<xref rid="figure2" ref-type="fig">Figure 2</xref>). Cluster 1 (n=12) was concentrated in the South and parts of the West, including Texas, Oklahoma, and several Rocky Mountain states. Cluster 2 (n=6) was scattered, appearing in parts of the central, Northeast, and Southeast. Cluster 3 (n=23), the largest cluster, covered much of the central and western regions, as well as parts of the northern region, including the Midwest, Great Plains, and Mountain West. Cluster 4 (n=11) was primarily located in the southern and southeastern US, with a few states in the West also assigned to this group.</p>
        <p>The policy characteristics of each cluster were assessed using factor scores (<xref ref-type="table" rid="table2">Table 2</xref> and Figure S3 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Cluster 1 had the highest scores across all 3 domains (Definition: 0.51, Reporting: 1.05, Alternative Response: 0.74), reflecting broad maltreatment definitions, strict mandated reporting policies with strong penalties, and extensive use of differential response systems. Cluster 2 showed the lowest Reporting score (–1.37), indicating fewer mandated reporting requirements and weaker penalties. Cluster 3 had the lowest Definition score (–0.43), suggesting narrower maltreatment definitions and more lenient safe haven policies. Cluster 4 had the lowest Alternative Response score (–1.05), reflecting the limited implementation of differential response systems and a likely greater reliance on traditional case-handling approaches.</p>
        <p>Regarding demographic and economic differences (<xref ref-type="table" rid="table3">Table 3</xref>), there were no significant differences in child population among clusters (cluster 1: 1.2 million; cluster 2: 2.1 million; cluster 3: 1.3 million; cluster 4: 1.4 million; <italic>P</italic>=.76). Demographically, there were minimal differences in age, gender, and race/ethnicity distributions across clusters. However, poverty status showed significant differences (Figure S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), with cluster 4 showing a higher proportion of children below poverty levels (cluster 1: 0.17; cluster 2: 0.16; cluster 3: 0.15; cluster 4: 0.21; <italic>P</italic>=.048).</p>
        <p>Regarding maltreatment outcomes (<xref ref-type="table" rid="table4">Table 4</xref>), cluster 2 exhibited the highest incidence of maltreatment, while cluster 1 had the lowest (cluster 1: mean 9.65, SD 4.10; cluster 2: mean 17.51, SD 6.13; cluster 3: mean 10.28, SD 6.47; cluster 4: mean 11.74, SD 6.47; <italic>P</italic>=.07). A similar trend was observed for maltreatment recurrence, with cluster 2 showing the highest rate (cluster 1: mean 0.52, SD 0.44; cluster 2: mean 1.74, SD 0.87; cluster 3: mean 0.66, SD 0.52; cluster 4: mean 0.68, SD 0.43; <italic>P</italic>&#60;.001). In contrast, there were no significant differences in maltreatment fatalities across clusters (cluster 1: mean 21.63, SD 14.83; cluster 2: mean 25.74, SD 14.70; cluster 3: mean 25.80, SD 15.64; cluster 4: mean 23.20, SD 18.80; <italic>P</italic>=.89). Pairwise comparisons (<xref rid="figure3" ref-type="fig">Figure 3</xref>) revealed that cluster 2 had a significantly higher incidence of maltreatment than Clusters 1 and 4 and a significantly higher recurrence of maltreatment than Clusters 1, 3, and 4.</p>
        <fig id="figure2" position="float">
          <label>Figure 2</label>
          <caption>
            <p>Geographic distribution of state child maltreatment policy clusters across 52 US jurisdictions based on 2021 State Child Abuse &#38; Neglect Policies Database data.</p>
          </caption>
          <graphic xlink:href="publichealth_v12i1e99643_fig2.png" alt-version="no" mimetype="image" position="float" xlink:type="simple"/>
        </fig>
        <table-wrap position="float" id="table2">
          <label>Table 2</label>
          <caption>
            <p>Comparison of latent child maltreatment policy factor scores across jurisdiction clusters identified through 2021 policy data from the State Child Abuse &#38; Neglect Policies Database.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="300"/>
            <col width="150"/>
            <col width="150"/>
            <col width="150"/>
            <col width="150"/>
            <col width="100"/>
            <thead>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Cluster 1 (n=12), mean (SD)</td>
                <td>Cluster 2 (n=6), mean (SD)</td>
                <td>Cluster 3 (n=23), mean (SD)</td>
                <td>Cluster 4 (n=11), mean (SD)</td>
                <td><italic>P</italic> value<sup>a</sup></td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Definition</td>
                <td>0.51 (1.19)</td>
                <td>0.00 (1.09)</td>
                <td>–0.43 (0.58)</td>
                <td>0.34 (1.15)</td>
                <td>.03</td>
              </tr>
              <tr valign="top">
                <td>Reporting</td>
                <td>1.05 (0.42)</td>
                <td>–1.37 (0.37)</td>
                <td>–0.53 (0.44)</td>
                <td>0.69 (0.83)</td>
                <td>&#60;.001</td>
              </tr>
              <tr valign="top">
                <td>Alternative response</td>
                <td>0.74 (0.76)</td>
                <td>–0.03 (0.79)</td>
                <td>0.12 (0.95)</td>
                <td>–1.05 (0.23)</td>
                <td>&#60;.001</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table2fn1">
              <p><sup>a</sup><italic>P</italic> values are based on 1-way ANOVA tests.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
        <table-wrap position="float" id="table3">
          <label>Table 3</label>
          <caption>
            <p>Comparison of demographic characteristics and poverty levels across state child maltreatment policy clusters among 52 US jurisdictions in 2021.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="30"/>
            <col width="0"/>
            <col width="0"/>
            <col width="270"/>
            <col width="0"/>
            <col width="150"/>
            <col width="0"/>
            <col width="150"/>
            <col width="0"/>
            <col width="150"/>
            <col width="0"/>
            <col width="150"/>
            <col width="0"/>
            <col width="100"/>
            <thead>
              <tr valign="top">
                <td colspan="5">
                  <break/>
                </td>
                <td colspan="2">Cluster 1 (n=12), mean (SD)</td>
                <td colspan="2">Cluster 2 (n=6), mean (SD)</td>
                <td colspan="2">Cluster 3 (n=23), mean (SD)</td>
                <td colspan="2">Cluster 4 (n=11), mean (SD)</td>
                <td><italic>P</italic> value<sup>a</sup></td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td colspan="5">Child population (in millions)</td>
                <td colspan="2">1.2 (2.0)</td>
                <td colspan="2">2.1 (1.2)</td>
                <td colspan="2">1.3 (1.7)</td>
                <td colspan="2">1.4 (1.1)</td>
                <td>.76</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Age (years)</td>
              </tr>
              <tr valign="top">
                <td colspan="2">
                  <break/>
                </td>
                <td colspan="2">&#60;5</td>
                <td colspan="2">0.25 (0.01)</td>
                <td colspan="2">0.25 (0.01)</td>
                <td colspan="2">0.26 (0.02)</td>
                <td colspan="2">0.25 (0.02)</td>
                <td colspan="2">.44</td>
              </tr>
              <tr valign="top">
                <td colspan="2">
                  <break/>
                </td>
                <td colspan="2">5-17</td>
                <td colspan="2">0.75 (0.01)</td>
                <td colspan="2">0.75 (0.01)</td>
                <td colspan="2">0.74 (0.02)</td>
                <td colspan="2">0.75 (0.02)</td>
                <td colspan="2">.44</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Gender</td>
              </tr>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td>Male</td>
                <td colspan="2">0.50 (0.01)</td>
                <td colspan="2">0.49 (0.01)</td>
                <td colspan="2">0.50 (0.01)</td>
                <td colspan="2">0.49 (0.01)</td>
                <td colspan="2">.13</td>
              </tr>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td>Female</td>
                <td colspan="2">0.50 (0.01)</td>
                <td colspan="2">0.51 (0.01)</td>
                <td colspan="2">0.50 (0.01)</td>
                <td colspan="2">0.51 (0.01)</td>
                <td colspan="2">.13</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Race and ethnicity</td>
              </tr>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td>White alone</td>
                <td colspan="2">0.62 (0.19)</td>
                <td colspan="2">0.64 (0.12)</td>
                <td colspan="2">0.70 (0.15)</td>
                <td colspan="2">0.67 (0.18)</td>
                <td colspan="2">.56</td>
              </tr>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td>Black alone</td>
                <td colspan="2">0.08 (0.07)</td>
                <td colspan="2">0.15 (0.10)</td>
                <td colspan="2">0.10 (0.11)</td>
                <td colspan="2">0.12 (0.11)</td>
                <td colspan="2">.52</td>
              </tr>
              <tr valign="top">
                <td colspan="3">
                  <break/>
                </td>
                <td>Hispanic or Latino</td>
                <td colspan="2">0.18 (0.14)</td>
                <td colspan="2">0.12 (0.06)</td>
                <td colspan="2">0.11 (0.10)</td>
                <td colspan="2">0.19 (0.28)</td>
                <td colspan="2">.46</td>
              </tr>
              <tr valign="top">
                <td colspan="14">Poverty status</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td colspan="4">Below 50% poverty level</td>
                <td colspan="2">0.06 (0.01)</td>
                <td colspan="2">0.06 (0.01)</td>
                <td colspan="2">0.06 (0.01)</td>
                <td colspan="2">0.08 (0.05)</td>
                <td>.11</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td colspan="4">Below poverty level</td>
                <td colspan="2">0.17 (0.04)</td>
                <td colspan="2">0.16 (0.02)</td>
                <td colspan="2">0.15 (0.03)</td>
                <td colspan="2">0.21 (0.10)</td>
                <td>.048</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td colspan="4">Below 125% poverty level</td>
                <td colspan="2">0.21 (0.05)</td>
                <td colspan="2">0.20 (0.03)</td>
                <td colspan="2">0.19 (0.03)</td>
                <td colspan="2">0.26 (0.12)</td>
                <td>.04</td>
              </tr>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td colspan="4">Below 400% poverty level</td>
                <td colspan="2">0.13 (0.03)</td>
                <td colspan="2">0.13 (0.02)</td>
                <td colspan="2">0.12 (0.02)</td>
                <td colspan="2">0.16 (0.09)</td>
                <td>.06</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table3fn1">
              <p><sup>a</sup><italic>P</italic> values are based on 1-way ANOVA tests.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
        <table-wrap position="float" id="table4">
          <label>Table 4</label>
          <caption>
            <p>Comparison of child maltreatment outcomes across state child maltreatment policy clusters among 52 US jurisdictions in 2021.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="300"/>
            <col width="150"/>
            <col width="150"/>
            <col width="150"/>
            <col width="150"/>
            <col width="100"/>
            <thead>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Cluster 1 (n=12), mean (SD)</td>
                <td>Cluster 2 (n=6), mean (SD)</td>
                <td>Cluster 3 (n=23), mean (SD)</td>
                <td>Cluster 4 (n=11), mean (SD)</td>
                <td><italic>P</italic> value<sup>a</sup></td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Incidence of maltreatment</td>
                <td>9.65 (4.10)</td>
                <td>17.51 (6.13)</td>
                <td>10.28 (6.47)</td>
                <td>11.74 (6.47)</td>
                <td>.07</td>
              </tr>
              <tr valign="top">
                <td>Maltreatment recurrence</td>
                <td>0.52 (0.44)</td>
                <td>1.74 (0.87)</td>
                <td>0.66 (0.52)</td>
                <td>0.68 (0.43)</td>
                <td>&#60;.001</td>
              </tr>
              <tr valign="top">
                <td>Maltreatment fatalities</td>
                <td>21.63 (14.83)</td>
                <td>25.74 (14.70)</td>
                <td>25.80 (15.64)</td>
                <td>23.20 (18.80)</td>
                <td>.89</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table4fn1">
              <p><sup>a</sup><italic>P</italic> values are based on 1-way ANOVA tests.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
        <fig id="figure3" position="float">
          <label>Figure 3</label>
          <caption>
            <p>Pairwise comparisons of maltreatment outcomes across state child maltreatment policy clusters among 52 US jurisdictions in 2021.</p>
          </caption>
          <graphic xlink:href="publichealth_v12i1e99643_fig3.png" alt-version="no" mimetype="image" position="float" xlink:type="simple"/>
        </fig>
      </sec>
      <sec>
        <title>Sensitivity Analyses</title>
        <p>The leave-one-jurisdiction-out sensitivity analysis demonstrated high cluster stability (mean ARI 0.984, range 0.726-1.000), suggesting that the identified policy configurations were not driven by individual jurisdictions.</p>
        <p>A sensitivity analysis using all screened-in NCANDS maltreatment records in 2021, regardless of substantiation or indication status, showed similar findings. Cluster 2 continued to have the highest maltreatment incidence and recurrence, although only the difference between cluster 2 and cluster 3 remained statistically significant (Figure S5 and Table S5 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p>
      </sec>
      <sec>
        <title>Subanalysis Results</title>
        <p>An independent clustering analysis based solely on the 18 fatality-related policy items (Figure S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) identified 4 fatality-policy clusters, which are distinct from the 4 policy clusters reported in the primary analysis. Among these fatality-policy clusters, cluster 2 had the least clear definition for child fatalities and near-fatalities, with the lowest scores in “child fatalities definition includes injury,” “child fatalities definition includes death of children in foster care,” “child fatalities definition includes other,” “child near-fatalities definition includes general reference to condition/injury,” and “child near-fatalities definition includes specific injury or treatment” (Table S6 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p>
        <p>Regarding outcomes, there were no significant differences among fatality-related policy clusters in terms of incidence of maltreatment or maltreatment recurrence (<xref ref-type="table" rid="table5">Table 5</xref>). However, a significant difference was found in maltreatment fatalities (cluster 1: mean 22.64, SD 11.82; cluster 2: mean 47.91, SD 8.35; cluster 3: mean 25.63, SD 18.62; cluster 4: mean 19.03, SD 11.64; <italic>P</italic>=.03).</p>
        <p>Pairwise comparisons showed that cluster 2 had significantly higher maltreatment fatalities compared to Clusters 1, 3, and 4 (<xref rid="figure4" ref-type="fig">Figure 4</xref>). The fatality distribution for cluster 2 was relatively narrow and centered around a higher median, indicating consistently elevated fatality rates. In contrast, cluster 3 showed a wider and more skewed distribution with greater variability, while clusters 1 and 4 showed more moderate and tightly clustered fatality rates with lower medians.</p>
        <table-wrap position="float" id="table5">
          <label>Table 5</label>
          <caption>
            <p>Comparison of maltreatment outcomes across jurisdictions grouped by fatality-related policy clusters.</p>
          </caption>
          <table width="1000" cellpadding="5" cellspacing="0" border="1" rules="groups" frame="hsides">
            <col width="280"/>
            <col width="150"/>
            <col width="140"/>
            <col width="140"/>
            <col width="140"/>
            <col width="150"/>
            <thead>
              <tr valign="top">
                <td>
                  <break/>
                </td>
                <td>Cluster 1 (n=15), mean (SD)</td>
                <td>Cluster 2 (n=3), mean (SD)</td>
                <td>Cluster 3 (n=20), mean (SD)</td>
                <td>Cluster 4 (n=14), mean (SD)</td>
                <td><italic>P</italic> value<sup>a</sup></td>
              </tr>
            </thead>
            <tbody>
              <tr valign="top">
                <td>Incidence of maltreatment</td>
                <td>10.57 (7.11)</td>
                <td>14.56 (6.21)</td>
                <td>11.21 (5.63)</td>
                <td>11.43 (7.35)</td>
                <td>.82</td>
              </tr>
              <tr valign="top">
                <td>Maltreatment recurrence</td>
                <td>0.61 (0.55)</td>
                <td>0.89 (0.21)</td>
                <td>0.75 (0.57)</td>
                <td>0.89 (0.84)</td>
                <td>.68</td>
              </tr>
              <tr valign="top">
                <td>Maltreatment fatalities</td>
                <td>22.64 (11.82)</td>
                <td>47.91 (8.35)</td>
                <td>25.63 (18.62)</td>
                <td>19.03 (11.64)</td>
                <td>.03</td>
              </tr>
            </tbody>
          </table>
          <table-wrap-foot>
            <fn id="table5fn1">
              <p><sup>a</sup><italic>P</italic> values are based on 1-way ANOVA tests.</p>
            </fn>
          </table-wrap-foot>
        </table-wrap>
        <fig id="figure4" position="float">
          <label>Figure 4</label>
          <caption>
            <p>Pairwise outcome differences for subanalysis based on the 18 maltreatment fatality policy items.</p>
          </caption>
          <graphic xlink:href="publichealth_v12i1e99643_fig4.png" alt-version="no" mimetype="image" position="float" xlink:type="simple"/>
        </fig>
      </sec>
    </sec>
    <sec sec-type="discussion">
      <title>Discussion</title>
      <sec>
        <title>Summary of Findings</title>
        <p>This study analyzed state child maltreatment policies and identified distinct policy clusters associated with maltreatment outcomes. Key findings indicated that states with the fewest mandated reporting requirements and weakest penalties (cluster 2) were associated with the highest substantiated or indicated maltreatment incidence and recurrence rates recorded in NCANDS, whereas states with the strongest reporting requirements, broadest maltreatment definitions, and most extensive use of alternative response (cluster 1) exhibited the lowest rates of substantiated or indicated maltreatment and recurrence. Additionally, states with broad maltreatment definitions and stringent reporting policies (cluster 4) had a high proportion of individuals below the poverty level, although maltreatment outcomes in these states were not significantly different.</p>
      </sec>
      <sec>
        <title>Interpretation of Findings</title>
        <p>Previous studies found that mandated reporting policies increased the number of maltreatment reports but did not have a significant impact on substantiated or indicated maltreatment incidence [<xref ref-type="bibr" rid="ref7">7</xref>,<xref ref-type="bibr" rid="ref54">54</xref>]. This study expanded upon them by suggesting that mandated reporting policies, when coupled with broad definitions of maltreatment and alternative response systems, were associated with lower rates of substantiated or indicated maltreatment and recurrence. These findings highlight the importance of considering the broader policy context in evaluating the role of mandated reporting.</p>
        <p>While previous studies had linked child poverty to higher maltreatment reports [<xref ref-type="bibr" rid="ref55">55</xref>,<xref ref-type="bibr" rid="ref56">56</xref>], this study found that clusters with broader definitions and limited alternative response systems tended to be more economically disadvantaged. Although maltreatment outcomes in these states were slightly higher, differences were not statistically significant. These findings suggest the importance of considering how policy frameworks may intersect with socioeconomic conditions to shape child welfare involvement.</p>
        <p>This relationship between child fatality definitions and maltreatment fatality rates has not been explicitly examined in prior research. These findings suggested that a lack of definitional specificity may influence reporting and intervention efforts, potentially contributing to higher fatality rates.</p>
      </sec>
      <sec>
        <title>Policy Implications</title>
        <p>The findings suggested several considerations for child maltreatment policy. First, states may benefit from integrating mandated reporting with coordinated prevention and response approaches. Strengthening reporting requirements while also implementing broad maltreatment definitions and alternative response systems was associated with lower rates of substantiated or indicated maltreatment and recurrence recorded in NCANDS. Developing comprehensive frameworks that balance reporting with both preventive services and supportive response options, such as alternative response systems, could help improve child welfare outcomes [<xref ref-type="bibr" rid="ref57">57</xref>-<xref ref-type="bibr" rid="ref59">59</xref>].</p>
        <p>Additionally, states with broad maltreatment definitions and stringent reporting policies had higher rates of poverty. While maltreatment outcomes in these states were not significantly different, the overlap between expansive definitions and economic disadvantage raises concerns that reports of maltreatment may, at times, reflect conditions of financial hardship rather than actual harm. This highlights the importance of refining definitions to ensure they do not conflate poverty with neglect [<xref ref-type="bibr" rid="ref60">60</xref>,<xref ref-type="bibr" rid="ref61">61</xref>].</p>
        <p>Standardizing child fatality definitions across jurisdictions could also enhance the accuracy of reporting and effectiveness of intervention efforts. Establishing clear and consistent criteria for classifying and reporting child fatalities related to maltreatment may improve cross-state comparisons and inform more effective prevention strategies [<xref ref-type="bibr" rid="ref62">62</xref>,<xref ref-type="bibr" rid="ref63">63</xref>].</p>
        <p>Lastly, the use of NLP and other advanced analytics in policy analysis demonstrated the potential for data-driven approaches in shaping more effective child welfare policies. Leveraging these technologies could provide valuable insights into policy development and evaluation.</p>
      </sec>
      <sec>
        <title>Limitations and Future Directions</title>
        <p>This study had several limitations. First, while the NLP models quantified policy characteristics, they may struggle with nuanced interpretations, particularly in complex policy language. Second, the study relied on state-reported NCANDS data, which may contain inconsistencies in reporting practices across jurisdictions. Third, the analysis was limited to policies in 2021, as more recent data were not available at the time of the study, potentially missing subsequent policy changes. Fourth, the cross-sectional design limited the ability to assess causal relationships between policy configurations and maltreatment outcomes. Finally, semantic alignment was challenging to evaluate due to the absence of a true gold standard, and reliance on proxy measures may not fully capture model accuracy.</p>
        <p>Beyond these limitations, several sources of structural uncertainty are worth noting, as their exact magnitude remains largely unknown. Cross-state variation in reporting practices means observed outcome differences may partly reflect how jurisdictions define and record cases rather than underlying policy effects alone. Lag times between policy enactment and implementation add further uncertainty, since coded 2021 policies may not yet have been fully operationalized, while outcomes that year may partly reflect earlier policy environments. Tracking inconsistencies across state data systems also introduces measurement noise that is difficult to quantify but may affect the stability of cluster-outcome associations.</p>
        <p>Future research should refine NLP approaches to enhance policy interpretation, explore longitudinal analyses to assess policy changes over time, and examine the impact of federal initiatives such as the Family First Prevention Services Act on state policy configurations and child welfare outcomes. With larger longitudinal datasets, future studies could also explore integrated machine learning or hierarchical modeling approaches to jointly model co-occurring policy configurations and multiple child welfare outcomes while accounting for temporal relationships and potential confounding factors. Recent advances in causal machine learning methods for policy evaluation [<xref ref-type="bibr" rid="ref64">64</xref>] could further help address confounding in future work assessing the causal impact of specific policy configurations on maltreatment outcomes. Finally, parametric approaches integrating multiple co-occurring policies to predict population-level outcomes have also been developed in other domains, such as epidemic modeling of nonpharmacological interventions [<xref ref-type="bibr" rid="ref65">65</xref>], and may offer a useful direction for future longitudinal extensions of this work.</p>
      </sec>
      <sec>
        <title>Conclusions</title>
        <p>This study identified distinct state-level child maltreatment policy configurations and examined their associations with maltreatment outcomes using an NLP-based analytical framework. This study advanced NLP-based analysis of child maltreatment policies across all US states, offering a scalable, reproducible approach to identify policy patterns. By integrating EFA and clustering, this study moved beyond descriptive comparisons and uncovered latent policy structures and their associations with demographic, economic, and child maltreatment outcomes. These findings offered a data-driven basis for comparing policies, identifying gaps, and informing reform.</p>
      </sec>
    </sec>
  </body>
  <back>
    <app-group>
      <supplementary-material id="app1">
        <label>Multimedia Appendix 1</label>
        <p>Supplementary figures and tables providing factor analysis diagnostics, cluster determinations, geographic distributions, outcome comparisons, and SCAN policy item evaluations.</p>
        <media xlink:href="publichealth_v12i1e99643_app1.docx" xlink:title="DOCX File , 260 KB"/>
      </supplementary-material>
    </app-group>
    <glossary>
      <title>Abbreviations</title>
      <def-list>
        <def-item>
          <term id="abb1">ARI</term>
          <def>
            <p>adjusted Rand index</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb2">BART</term>
          <def>
            <p>bidirectional and auto-regressive transformer</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb3">BERT</term>
          <def>
            <p>bidirectional encoder representations from transformers</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb4">BIC</term>
          <def>
            <p>Bayesian information criterion</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb5">BSS/WSS</term>
          <def>
            <p>between-cluster sum of squares to within-cluster sum of squares</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb6">CAPTA</term>
          <def>
            <p>Child Abuse Prevention and Treatment Act</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb7">CH</term>
          <def>
            <p>Calinski-Harabasz score</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb8">DeBERTa</term>
          <def>
            <p>decoding-enhanced BERT with disentangled attention</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb9">EFA</term>
          <def>
            <p>exploratory factor analysis</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb10">LLM</term>
          <def>
            <p>large language model</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb11">MRSF</term>
          <def>
            <p>multiple R-squared of scores with factors</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb12">NCANDS</term>
          <def>
            <p>National Child Abuse and Neglect Data System</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb13">NDACAN</term>
          <def>
            <p>National Data Archive on Child Abuse and Neglect</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb14">NLP</term>
          <def>
            <p>natural language processing</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb15">RoBERTa</term>
          <def>
            <p>robustly optimized BERT approach</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb16">SCAN</term>
          <def>
            <p>State Child Abuse &#38; Neglect Policies Database</p>
          </def>
        </def-item>
        <def-item>
          <term id="abb17">STROBE</term>
          <def>
            <p>Strengthening the Reporting of Observational Studies in Epidemiology</p>
          </def>
        </def-item>
      </def-list>
    </glossary>
    <ack>
      <p>Generative AI tools were not used in the generation, analysis, or interpretation of this manuscript. Grammarly was used for grammar and language editing assistance.</p>
    </ack>
    <notes>
      <sec>
        <title>Funding</title>
        <p>The authors declared no financial support was received for this work.</p>
      </sec>
    </notes>
    <notes>
      <sec>
        <title>Data Availability</title>
        <p>The datasets used in this study are available from the National Data Archive on Child Abuse and Neglect (NDACAN) but are not publicly available due to data use restrictions. Both the National Child Abuse and Neglect Data System (NCANDS) and the State Child Abuse &#38; Neglect Policies (SCAN) dataset require submission of a data use agreement and approval by NDACAN. Researchers may request access through NDACAN [<xref ref-type="bibr" rid="ref66">66</xref>] and must comply with all data security and use requirements. Demographic and economic data from the American Community Survey are publicly available through the US Census Bureau.</p>
      </sec>
    </notes>
    <fn-group>
      <fn fn-type="con">
        <p>Conceptualization: ZL, RAE</p>
        <p>Data curation: ZL</p>
        <p>Formal analysis: ZL, NS, LNM, RAE, NJ</p>
        <p>Investigation: ZL, NS, LNM, RAE, NJ</p>
        <p>Methodology: ZL, NS, LNM</p>
        <p>Project administration: ZL, RAE, NJ</p>
        <p>Resources: RAE, NJ</p>
        <p>Software: ZL</p>
        <p>Supervision: RAE, NJ, NS, LNM</p>
        <p>Validation: ZL, NS, LNM, RAE, NJ</p>
        <p>Visualization: ZL</p>
        <p>Writing—original draft: ZL</p>
        <p>Writing—review and editing: ZL, RAE, NS, LNM, NJ</p>
      </fn>
      <fn fn-type="conflict">
        <p>None declared.</p>
      </fn>
    </fn-group>
    <ref-list>
      <ref id="ref1">
        <label>1</label>
        <nlm-citation citation-type="web">
          <article-title>Child maltreatment</article-title>
          <source>Administration for Children and Families</source>
          <year>2025</year>
          <access-date>2026-09-25</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://acf.gov/sites/default/files/documents/cb/cm2023.pdf">https://acf.gov/sites/default/files/documents/cb/cm2023.pdf</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref2">
        <label>2</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Norman</surname>
              <given-names>RE</given-names>
            </name>
            <name name-style="western">
              <surname>Byambaa</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>De</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Butchart</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Scott</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Vos</surname>
              <given-names>T</given-names>
            </name>
          </person-group>
          <article-title>The long-term health consequences of child physical abuse, emotional abuse, and neglect: a systematic review and meta-analysis</article-title>
          <source>PLoS Med</source>
          <year>2012</year>
          <volume>9</volume>
          <issue>11</issue>
          <fpage>e1001349</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://dx.plos.org/10.1371/journal.pmed.1001349"/>
          </comment>
          <pub-id pub-id-type="doi">10.1371/journal.pmed.1001349</pub-id>
          <pub-id pub-id-type="medline">23209385</pub-id>
          <pub-id pub-id-type="pii">PMEDICINE-D-12-00821</pub-id>
          <pub-id pub-id-type="pmcid">PMC3507962</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref3">
        <label>3</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Merrick</surname>
              <given-names>MT</given-names>
            </name>
            <name name-style="western">
              <surname>Ford</surname>
              <given-names>DC</given-names>
            </name>
            <name name-style="western">
              <surname>Ports</surname>
              <given-names>KA</given-names>
            </name>
            <name name-style="western">
              <surname>Guinn</surname>
              <given-names>AS</given-names>
            </name>
          </person-group>
          <article-title>Prevalence of adverse childhood experiences from the 2011-2014 behavioral risk factor surveillance system in 23 states</article-title>
          <source>JAMA Pediatr</source>
          <year>2018</year>
          <volume>172</volume>
          <issue>11</issue>
          <fpage>1038</fpage>
          <lpage>1044</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/30242348"/>
          </comment>
          <pub-id pub-id-type="doi">10.1001/jamapediatrics.2018.2537</pub-id>
          <pub-id pub-id-type="medline">30242348</pub-id>
          <pub-id pub-id-type="pii">2702204</pub-id>
          <pub-id pub-id-type="pmcid">PMC6248156</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref4">
        <label>4</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Danese</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Widom</surname>
              <given-names>CS</given-names>
            </name>
          </person-group>
          <article-title>Objective and subjective experiences of child maltreatment and their relationships with psychopathology</article-title>
          <source>Nat Hum Behav</source>
          <year>2020</year>
          <volume>4</volume>
          <issue>8</issue>
          <fpage>811</fpage>
          <lpage>818</lpage>
          <pub-id pub-id-type="doi">10.1038/s41562-020-0880-3</pub-id>
          <pub-id pub-id-type="medline">32424258</pub-id>
          <pub-id pub-id-type="pii">10.1038/s41562-020-0880-3</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref5">
        <label>5</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Petruccelli</surname>
              <given-names>K</given-names>
            </name>
            <name name-style="western">
              <surname>Davis</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Berman</surname>
              <given-names>T</given-names>
            </name>
          </person-group>
          <article-title>Adverse childhood experiences and associated health outcomes: a systematic review and meta-analysis</article-title>
          <source>Child Abuse Negl</source>
          <year>2019</year>
          <volume>97</volume>
          <fpage>104127</fpage>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2019.104127</pub-id>
          <pub-id pub-id-type="medline">31454589</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(19)30304-7</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref6">
        <label>6</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Klika</surname>
              <given-names>JB</given-names>
            </name>
            <name name-style="western">
              <surname>Rosenzweig</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Merrick</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Economic burden of known cases of child maltreatment from 2018 in each state</article-title>
          <source>Child Adolesc Soc Work J</source>
          <year>2020</year>
          <volume>37</volume>
          <issue>3</issue>
          <fpage>227</fpage>
          <lpage>234</lpage>
          <pub-id pub-id-type="doi">10.1007/s10560-020-00665-5</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref7">
        <label>7</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Day</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Tach</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Mihalec-Adkins</surname>
              <given-names>B</given-names>
            </name>
          </person-group>
          <article-title>State child welfare policies and the measurement of child maltreatment in the United States</article-title>
          <source>Child Maltreat</source>
          <year>2022</year>
          <volume>27</volume>
          <issue>3</issue>
          <fpage>411</fpage>
          <lpage>422</lpage>
          <pub-id pub-id-type="doi">10.1177/10775595211006464</pub-id>
          <pub-id pub-id-type="medline">33832331</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref8">
        <label>8</label>
        <nlm-citation citation-type="web">
          <article-title>Mandated reporting</article-title>
          <source>Child Welfare Information Gateway</source>
          <access-date>2026-04-16</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.childwelfare.gov/topics/safety-and-risk/mandated-reporting/?top=78">https://www.childwelfare.gov/topics/safety-and-risk/mandated-reporting/?top=78</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref9">
        <label>9</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Johnson-Motoyama</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Ginther</surname>
              <given-names>DK</given-names>
            </name>
            <name name-style="western">
              <surname>Phillips</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Beer</surname>
              <given-names>OWJ</given-names>
            </name>
            <name name-style="western">
              <surname>Merkel-Holguin</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Fluke</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>Differential response and the reduction of child maltreatment and foster care services utilization in the U.S. from 2004 to 2017</article-title>
          <source>Child Maltreat</source>
          <year>2023</year>
          <volume>28</volume>
          <issue>1</issue>
          <fpage>152</fpage>
          <lpage>162</lpage>
          <pub-id pub-id-type="doi">10.1177/10775595211065761</pub-id>
          <pub-id pub-id-type="medline">35062827</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref10">
        <label>10</label>
        <nlm-citation citation-type="web">
          <article-title>The Child Abuse Prevention and Treatment Act (CAPTA): background, programs, and funding</article-title>
          <source>EveryCSReport</source>
          <year>2009</year>
          <access-date>2026-04-16</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.everycrsreport.com/reports/R40899.html">https://www.everycrsreport.com/reports/R40899.html</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref11">
        <label>11</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>LaBrenz</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Kim</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Baiden</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Shipe</surname>
              <given-names>SL</given-names>
            </name>
            <name name-style="western">
              <surname>Littleton</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Choi</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Bai</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Stargel</surname>
              <given-names>L</given-names>
            </name>
          </person-group>
          <article-title>State child maltreatment policies and disparities in substantiation: a study of state-administered child welfare systems in the U.S</article-title>
          <source>Child Maltreat</source>
          <year>2023</year>
          <volume>28</volume>
          <issue>4</issue>
          <fpage>700</fpage>
          <lpage>712</lpage>
          <pub-id pub-id-type="doi">10.1177/10775595221143136</pub-id>
          <pub-id pub-id-type="medline">36458462</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref12">
        <label>12</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Puls</surname>
              <given-names>HT</given-names>
            </name>
            <name name-style="western">
              <surname>Hall</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Anderst</surname>
              <given-names>JD</given-names>
            </name>
            <name name-style="western">
              <surname>Gurley</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Perrin</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Chung</surname>
              <given-names>PJ</given-names>
            </name>
          </person-group>
          <article-title>State spending on public benefit programs and child maltreatment</article-title>
          <source>Pediatrics</source>
          <year>2021</year>
          <volume>148</volume>
          <issue>5</issue>
          <fpage>5</fpage>
          <pub-id pub-id-type="doi">10.1542/peds.2021-050685</pub-id>
          <pub-id pub-id-type="medline">34663680</pub-id>
          <pub-id pub-id-type="pii">peds.2021-050685</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref13">
        <label>13</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Gunes</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Florczak</surname>
              <given-names>CK</given-names>
            </name>
          </person-group>
          <article-title>Multiclass classification of policy documents with large language models</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on October 12, 2023</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://arxiv.org/abs/2310.08167v1"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref14">
        <label>14</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Osnabrügge</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Vannoni</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Quality of legislation and compliance: a natural language processing approach</article-title>
          <source>PSRM</source>
          <year>2024</year>
          <volume>13</volume>
          <issue>3</issue>
          <fpage>736</fpage>
          <lpage>744</lpage>
          <pub-id pub-id-type="doi">10.1017/psrm.2024.21</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref15">
        <label>15</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Swarnakar</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Modi</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>NLP for climate policy: creating a knowledge platform for holistic and effective climate action</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on May 12, 2021</comment>
          <pub-id pub-id-type="doi">10.48550/arXiv.2105.05621</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref16">
        <label>16</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Planas</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Firebanks-Quevedo</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Naydenova</surname>
              <given-names>G</given-names>
            </name>
            <name name-style="western">
              <surname>Sharma</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Taylor</surname>
              <given-names>C</given-names>
            </name>
            <name name-style="western">
              <surname>Buckingham</surname>
              <given-names>K</given-names>
            </name>
          </person-group>
          <article-title>Beyond modeling: NLP pipeline for efficient environmental policy analysis</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on January 8, 2022</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/2201.07105"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref17">
        <label>17</label>
        <nlm-citation citation-type="web">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Nay</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>Natural language processing and machine learning for law and policy texts</article-title>
          <source>SSRN</source>
          <year>2018</year>
          <access-date>2024-08-22</access-date>
          <publisher-loc>Rochester, NY</publisher-loc>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://papers.ssrn.com/abstract=3438276">https://papers.ssrn.com/abstract=3438276</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref18">
        <label>18</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Rosholm</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Bodilsen</surname>
              <given-names>ST</given-names>
            </name>
            <name name-style="western">
              <surname>Michel</surname>
              <given-names>B</given-names>
            </name>
            <name name-style="western">
              <surname>Nielsen</surname>
              <given-names>AS</given-names>
            </name>
          </person-group>
          <article-title>Predictive risk modeling for child maltreatment detection and enhanced decision-making: evidence from Danish administrative data</article-title>
          <source>PLoS One</source>
          <year>2024</year>
          <volume>19</volume>
          <issue>7</issue>
          <fpage>e0305974</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://dx.plos.org/10.1371/journal.pone.0305974"/>
          </comment>
          <pub-id pub-id-type="doi">10.1371/journal.pone.0305974</pub-id>
          <pub-id pub-id-type="medline">38985689</pub-id>
          <pub-id pub-id-type="pii">PONE-D-23-30480</pub-id>
          <pub-id pub-id-type="pmcid">PMC11236184</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref19">
        <label>19</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Ahn</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>An</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Jonson-Reid</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Palmer</surname>
              <given-names>L</given-names>
            </name>
          </person-group>
          <article-title>Leveraging machine learning for effective child maltreatment prevention: a case study of home visiting service assessments</article-title>
          <source>Child Abuse Negl</source>
          <year>2024</year>
          <volume>151</volume>
          <fpage>106706</fpage>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2024.106706</pub-id>
          <pub-id pub-id-type="medline">38428267</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(24)00089-9</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref20">
        <label>20</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Negriff</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Lynch</surname>
              <given-names>FL</given-names>
            </name>
            <name name-style="western">
              <surname>Cronkite</surname>
              <given-names>DJ</given-names>
            </name>
            <name name-style="western">
              <surname>Pardee</surname>
              <given-names>RE</given-names>
            </name>
            <name name-style="western">
              <surname>Penfold</surname>
              <given-names>RB</given-names>
            </name>
          </person-group>
          <article-title>Using natural language processing to identify child maltreatment in health systems</article-title>
          <source>Child Abuse Negl</source>
          <year>2023</year>
          <volume>138</volume>
          <fpage>106090</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/36758373"/>
          </comment>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2023.106090</pub-id>
          <pub-id pub-id-type="medline">36758373</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(23)00071-6</pub-id>
          <pub-id pub-id-type="pmcid">PMC9984187</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref21">
        <label>21</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Annapragada</surname>
              <given-names>AV</given-names>
            </name>
            <name name-style="western">
              <surname>Donaruma-Kwoh</surname>
              <given-names>MM</given-names>
            </name>
            <name name-style="western">
              <surname>Annapragada</surname>
              <given-names>AV</given-names>
            </name>
            <name name-style="western">
              <surname>Starosolski</surname>
              <given-names>ZA</given-names>
            </name>
          </person-group>
          <article-title>A natural language processing and deep learning approach to identify child abuse from pediatric electronic medical records</article-title>
          <source>PLoS One</source>
          <year>2021</year>
          <volume>16</volume>
          <issue>2</issue>
          <fpage>e0247404</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://dx.plos.org/10.1371/journal.pone.0247404"/>
          </comment>
          <pub-id pub-id-type="doi">10.1371/journal.pone.0247404</pub-id>
          <pub-id pub-id-type="medline">33635890</pub-id>
          <pub-id pub-id-type="pii">PONE-D-20-16162</pub-id>
          <pub-id pub-id-type="pmcid">PMC7909689</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref22">
        <label>22</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Victor</surname>
              <given-names>BG</given-names>
            </name>
            <name name-style="western">
              <surname>Perron</surname>
              <given-names>BE</given-names>
            </name>
            <name name-style="western">
              <surname>Sokol</surname>
              <given-names>RL</given-names>
            </name>
            <name name-style="western">
              <surname>Fedina</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Ryan</surname>
              <given-names>JP</given-names>
            </name>
          </person-group>
          <article-title>Automated identification of domestic violence in written child welfare records: leveraging text mining and machine learning to enhance social work research and evaluation</article-title>
          <source>J Soc Social Work Res</source>
          <year>2021</year>
          <volume>12</volume>
          <issue>4</issue>
          <fpage>631</fpage>
          <lpage>655</lpage>
          <pub-id pub-id-type="doi">10.1086/712734</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref23">
        <label>23</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Niu</surname>
              <given-names>F</given-names>
            </name>
            <name name-style="western">
              <surname>Dhamayanti</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Setiawati</surname>
              <given-names>EP</given-names>
            </name>
            <name name-style="western">
              <surname>Arisanti</surname>
              <given-names>N</given-names>
            </name>
          </person-group>
          <article-title>Artificial intelligence for early detection of child maltreatment in healthcare: a narrative review integrating technical, ethical, clinical, and governance perspectives</article-title>
          <source>Children and Youth Services Review</source>
          <year>2026</year>
          <volume>188</volume>
          <fpage>109106</fpage>
          <pub-id pub-id-type="doi">10.1016/j.childyouth.2026.109106</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref24">
        <label>24</label>
        <nlm-citation citation-type="web">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Fox-Sowell</surname>
              <given-names>S</given-names>
            </name>
          </person-group>
          <article-title>Pennsylvania county taps natural language processing to help child welfare caseworkers</article-title>
          <source>StateScoop</source>
          <year>2023</year>
          <access-date>2025-02-22</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://statescoop.com/nlp-ai-washington-county-pennsylvania-child-welfare-caseworkers/">https://statescoop.com/nlp-ai-washington-county-pennsylvania-child-welfare-caseworkers/</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref25">
        <label>25</label>
        <nlm-citation citation-type="confproc">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Saxena</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Moon</surname>
              <given-names>ESY</given-names>
            </name>
            <name name-style="western">
              <surname>Chaurasia</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Guan</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Guha</surname>
              <given-names>S</given-names>
            </name>
          </person-group>
          <article-title>Rethinking 'Risk' in algorithmic systems through a computational narrative analysis of casenotes in child-welfare</article-title>
          <year>2023</year>
          <conf-name>Proceedings of the 2023 CHI Conference on Human Factors in Computing Systems</conf-name>
          <conf-date>2026 July 23</conf-date>
          <conf-loc>New York, NY, USA</conf-loc>
          <publisher-name>Association for Computing Machinery</publisher-name>
          <fpage>1</fpage>
          <lpage>19</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://dl.acm.org/doi/10.1145/3544548.3581308"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref26">
        <label>26</label>
        <nlm-citation citation-type="web">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Weigensberg</surname>
              <given-names>EC</given-names>
            </name>
            <name name-style="western">
              <surname>Islam</surname>
              <given-names>N</given-names>
            </name>
            <name name-style="western">
              <surname>Knab</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Grider</surname>
              <given-names>MA</given-names>
            </name>
            <name name-style="western">
              <surname>Page</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Larson</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>State Child Abuse and Neglect (SCAN) Policies Database 2019-2021</article-title>
          <source>National Data Archive on Child Abuse and Neglect</source>
          <year>2022</year>
          <access-date>2024-02-20</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.ndacan.acf.hhs.gov/datasets/dataset-details.cfm?ID=268">https://www.ndacan.acf.hhs.gov/datasets/dataset-details.cfm?ID=268</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref27">
        <label>27</label>
        <nlm-citation citation-type="web">
          <article-title>National Child Abuse and Neglect Data System (NCANDS) Agency File</article-title>
          <source>National Data Archive on Child Abuse and Neglect</source>
          <year>2024</year>
          <access-date>2024-08-08</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.ndacan.acf.hhs.gov/datasets/datasets-list-ncands-state-agency-file.cfm">https://www.ndacan.acf.hhs.gov/datasets/datasets-list-ncands-state-agency-file.cfm</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref28">
        <label>28</label>
        <nlm-citation citation-type="web">
          <article-title>National Child Abuse and Neglect Data System (NCANDS) Child File</article-title>
          <source>National Data Archive on Child Abuse and Neglect</source>
          <year>2024</year>
          <access-date>2024-08-08</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.ndacan.acf.hhs.gov/datasets/datasets-list-ncands-child-file.cfm">https://www.ndacan.acf.hhs.gov/datasets/datasets-list-ncands-child-file.cfm</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref29">
        <label>29</label>
        <nlm-citation citation-type="web">
          <article-title>American Community Survey</article-title>
          <source>Census.gov</source>
          <year>2024</year>
          <access-date>2025-02-25</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.census.gov/programs-surveys/acs">https://www.census.gov/programs-surveys/acs</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref30">
        <label>30</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Yin</surname>
              <given-names>W</given-names>
            </name>
            <name name-style="western">
              <surname>Hay</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Roth</surname>
              <given-names>D</given-names>
            </name>
          </person-group>
          <article-title>Benchmarking zero-shot text classification: datasets, evaluation and entailment approach</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on August 31, 2019</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/1909.00161"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref31">
        <label>31</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Lewis</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Goyal</surname>
              <given-names>N</given-names>
            </name>
            <name name-style="western">
              <surname>Ghazvininejad</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Mohamed</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Levy</surname>
              <given-names>O</given-names>
            </name>
          </person-group>
          <article-title>BART: denoising sequence-to-sequence pre-training for natural language generation, translation, and comprehension</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on October 29, 2019</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/1910.13461"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref32">
        <label>32</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Devlin</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Chang</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Lee</surname>
              <given-names>K</given-names>
            </name>
            <name name-style="western">
              <surname>Toutanova</surname>
              <given-names>K</given-names>
            </name>
          </person-group>
          <article-title>BERT: pre-training of deep bidirectional transformers for language understanding</article-title>
          <source>arXiv</source>
          <comment> Preprint posted online on October 11, 2018</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/1810.04805"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref33">
        <label>33</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Ott</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Goyal</surname>
              <given-names>N</given-names>
            </name>
            <name name-style="western">
              <surname>Du</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Joshi</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Chen</surname>
              <given-names>D</given-names>
            </name>
          </person-group>
          <article-title>RoBERTa: a robustly optimized BERT pretraining approach</article-title>
          <source>arXiv</source>
          <comment>Preprint posted online on July 26, 2019</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/1907.11692"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref34">
        <label>34</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>He</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>X</given-names>
            </name>
            <name name-style="western">
              <surname>Gao</surname>
              <given-names>J</given-names>
            </name>
            <name name-style="western">
              <surname>Chen</surname>
              <given-names>W</given-names>
            </name>
          </person-group>
          <article-title>DeBERTa: decoding-enhanced BERT with disentangled attention</article-title>
          <source>arXiv</source>
          <comment> Preprint posted online on June 5, 2020</comment>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="http://arxiv.org/abs/2006.03654"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref35">
        <label>35</label>
        <nlm-citation citation-type="web">
          <article-title>Microsoft 365 Copilot</article-title>
          <source>Microsoft</source>
          <year>2024</year>
          <access-date>2025-02-28</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://copilot.microsoft.com">https://copilot.microsoft.com</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref36">
        <label>36</label>
        <nlm-citation citation-type="web">
          <article-title>Meta Llama 3.1</article-title>
          <source>Meta AI</source>
          <year>2024</year>
          <access-date>2024-08-16</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://ai.meta.com/blog/meta-llama-3-1/">https://ai.meta.com/blog/meta-llama-3-1/</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref37">
        <label>37</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>B</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Chen</surname>
              <given-names>F</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Kuo</surname>
              <given-names>CCJ</given-names>
            </name>
          </person-group>
          <article-title>Evaluating word embedding models: methods and experimental results</article-title>
          <source>SIP</source>
          <year>2019</year>
          <volume>8</volume>
          <issue>1</issue>
          <pub-id pub-id-type="doi">10.1017/atsip.2019.12</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref38">
        <label>38</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Qiu</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Li</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Li</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Jiang</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Hu</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Yang</surname>
              <given-names>L</given-names>
            </name>
          </person-group>
          <person-group person-group-type="editor">
            <name name-style="western">
              <surname>Sun</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Wang</surname>
              <given-names>X</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>Z</given-names>
            </name>
            <name name-style="western">
              <surname>Liu</surname>
              <given-names>Y</given-names>
            </name>
          </person-group>
          <article-title>Revisiting correlations between intrinsic and extrinsic evaluations of word embeddings</article-title>
          <source>Chinese Computational Linguistics and Natural Language Processing Based on Naturally Annotated Big Data</source>
          <year>2018</year>
          <publisher-loc>Cham</publisher-loc>
          <publisher-name>Springer International Publishing</publisher-name>
          <fpage>209</fpage>
          <lpage>921</lpage>
        </nlm-citation>
      </ref>
      <ref id="ref39">
        <label>39</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Gorsuch</surname>
              <given-names>RL</given-names>
            </name>
          </person-group>
          <source>Factor Analysis. 2nd ed</source>
          <year>1983</year>
          <publisher-loc>New York</publisher-loc>
          <publisher-name>Psychology Press</publisher-name>
          <fpage>448</fpage>
        </nlm-citation>
      </ref>
      <ref id="ref40">
        <label>40</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Hartigan</surname>
              <given-names>JA</given-names>
            </name>
            <name name-style="western">
              <surname>Wong</surname>
              <given-names>MA</given-names>
            </name>
          </person-group>
          <article-title>Algorithm AS 136: A k-means clustering algorithm</article-title>
          <source>J R Stat Soc Series C Appl Stat</source>
          <year>1979</year>
          <volume>28</volume>
          <issue>1</issue>
          <fpage>100</fpage>
          <lpage>108</lpage>
          <pub-id pub-id-type="doi">10.2307/2346830</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref41">
        <label>41</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Sakamoto</surname>
              <given-names>Y</given-names>
            </name>
            <name name-style="western">
              <surname>Ishiguro</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Kitagawa</surname>
              <given-names>G</given-names>
            </name>
          </person-group>
          <source>Akaike Information Criterion Statistics</source>
          <year>1986</year>
          <publisher-loc>Netherland</publisher-loc>
          <publisher-name>Springer</publisher-name>
        </nlm-citation>
      </ref>
      <ref id="ref42">
        <label>42</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Calinski</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Harabasz</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>A dendrite method for cluster analysis</article-title>
          <source>Commun Stat Simul Comput</source>
          <year>1974</year>
          <volume>3</volume>
          <issue>1</issue>
          <fpage>1</fpage>
          <lpage>27</lpage>
          <pub-id pub-id-type="doi">10.1080/03610917408548446</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref43">
        <label>43</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Kriegel</surname>
              <given-names>HP</given-names>
            </name>
            <name name-style="western">
              <surname>Schubert</surname>
              <given-names>E</given-names>
            </name>
            <name name-style="western">
              <surname>Zimek</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>The (black) art of runtime evaluation: are we comparing algorithms or implementations?</article-title>
          <source>Knowl Inf Syst</source>
          <year>2016</year>
          <volume>52</volume>
          <issue>2</issue>
          <fpage>341</fpage>
          <lpage>378</lpage>
          <pub-id pub-id-type="doi">10.1007/s10115-016-1004-2</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref44">
        <label>44</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Cattell</surname>
              <given-names>RB</given-names>
            </name>
          </person-group>
          <article-title>The scree test for the number of factors</article-title>
          <source>Multivariate Behav Res</source>
          <year>1966</year>
          <volume>1</volume>
          <issue>2</issue>
          <fpage>245</fpage>
          <lpage>276</lpage>
          <pub-id pub-id-type="doi">10.1207/s15327906mbr0102_10</pub-id>
          <pub-id pub-id-type="medline">26828106</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref45">
        <label>45</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Kaiser</surname>
              <given-names>HF</given-names>
            </name>
          </person-group>
          <article-title>A second generation little jiffy</article-title>
          <source>Psychometrika</source>
          <year>2025</year>
          <volume>35</volume>
          <issue>4</issue>
          <fpage>401</fpage>
          <lpage>415</lpage>
          <pub-id pub-id-type="doi">10.1007/BF02291817</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref46">
        <label>46</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Grice</surname>
              <given-names>JW</given-names>
            </name>
          </person-group>
          <article-title>Computing and evaluating factor scores</article-title>
          <source>Psychol Methods</source>
          <year>2001</year>
          <volume>6</volume>
          <issue>4</issue>
          <fpage>430</fpage>
          <lpage>450</lpage>
          <pub-id pub-id-type="medline">11778682</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref47">
        <label>47</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>McDonald</surname>
              <given-names>RP</given-names>
            </name>
          </person-group>
          <source>Test Theory: A Unified Treatment</source>
          <year>1999</year>
          <publisher-loc>New Jersey, USA</publisher-loc>
          <publisher-name>Lawrence Erlbaum Associates Publishers</publisher-name>
        </nlm-citation>
      </ref>
      <ref id="ref48">
        <label>48</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Rousseeuw</surname>
              <given-names>PJ</given-names>
            </name>
          </person-group>
          <article-title>Silhouettes: a graphical aid to the interpretation and validation of cluster analysis</article-title>
          <source>J Comput Appl Math</source>
          <year>1987</year>
          <volume>20</volume>
          <fpage>53</fpage>
          <lpage>65</lpage>
          <pub-id pub-id-type="doi">10.1016/0377-0427(87)90125-7</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref49">
        <label>49</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Fisher</surname>
              <given-names>RA</given-names>
            </name>
          </person-group>
          <person-group person-group-type="editor">
            <name name-style="western">
              <surname>Kotz</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Johnson</surname>
              <given-names>NL</given-names>
            </name>
          </person-group>
          <article-title>Statistical methods for research workers</article-title>
          <source>Breakthroughs in Statistics: Methodology and Distribution</source>
          <year>1992</year>
          <publisher-loc>New York, NY</publisher-loc>
          <publisher-name>Springer</publisher-name>
          <fpage>66</fpage>
          <lpage>70</lpage>
        </nlm-citation>
      </ref>
      <ref id="ref50">
        <label>50</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Pearson</surname>
              <given-names>KX</given-names>
            </name>
          </person-group>
          <article-title>On the criterion that a given system of deviations from the probable in the case of a correlated system of variables is such that it can be reasonably supposed to have arisen from random sampling</article-title>
          <source>Lond Edinb Dubl Philos Mag</source>
          <year>2009</year>
          <volume>50</volume>
          <issue>302</issue>
          <fpage>157</fpage>
          <lpage>175</lpage>
          <pub-id pub-id-type="doi">10.1080/14786440009463897</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref51">
        <label>51</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Mudge</surname>
              <given-names>JF</given-names>
            </name>
            <name name-style="western">
              <surname>Baker</surname>
              <given-names>LF</given-names>
            </name>
            <name name-style="western">
              <surname>Edge</surname>
              <given-names>CB</given-names>
            </name>
            <name name-style="western">
              <surname>Houlahan</surname>
              <given-names>JE</given-names>
            </name>
          </person-group>
          <article-title>Setting an optimal α that minimizes errors in null hypothesis significance tests</article-title>
          <source>PLoS One</source>
          <year>2012</year>
          <volume>7</volume>
          <issue>2</issue>
          <fpage>e32734</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://dx.plos.org/10.1371/journal.pone.0032734"/>
          </comment>
          <pub-id pub-id-type="doi">10.1371/journal.pone.0032734</pub-id>
          <pub-id pub-id-type="medline">22389720</pub-id>
          <pub-id pub-id-type="pii">PONE-D-11-11244</pub-id>
          <pub-id pub-id-type="pmcid">PMC3289673</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref52">
        <label>52</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Cohen</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <source>Statistical Power Analysis for the Behavioral Sciences. 2nd ed</source>
          <year>1988</year>
          <publisher-loc>New York</publisher-loc>
          <publisher-name>Routledge</publisher-name>
          <fpage>567</fpage>
        </nlm-citation>
      </ref>
      <ref id="ref53">
        <label>53</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Hubert</surname>
              <given-names>L</given-names>
            </name>
            <name name-style="western">
              <surname>Arabie</surname>
              <given-names>P</given-names>
            </name>
          </person-group>
          <article-title>Comparing partitions</article-title>
          <source>J Classif</source>
          <year>1985</year>
          <volume>2</volume>
          <issue>1</issue>
          <fpage>193</fpage>
          <lpage>218</lpage>
          <pub-id pub-id-type="doi">10.1007/bf01908075</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref54">
        <label>54</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Ho</surname>
              <given-names>GWK</given-names>
            </name>
            <name name-style="western">
              <surname>Gross</surname>
              <given-names>DA</given-names>
            </name>
            <name name-style="western">
              <surname>Bettencourt</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Universal mandatory reporting policies and the odds of identifying child physical abuse</article-title>
          <source>Am J Public Health</source>
          <year>2017</year>
          <volume>107</volume>
          <issue>5</issue>
          <fpage>709</fpage>
          <lpage>716</lpage>
          <pub-id pub-id-type="doi">10.2105/AJPH.2017.303667</pub-id>
          <pub-id pub-id-type="medline">28323475</pub-id>
          <pub-id pub-id-type="pmcid">PMC5388942</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref55">
        <label>55</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Kim</surname>
              <given-names>H</given-names>
            </name>
            <name name-style="western">
              <surname>Drake</surname>
              <given-names>B</given-names>
            </name>
          </person-group>
          <article-title>Has the relationship between community poverty and child maltreatment report rates become stronger or weaker over time?</article-title>
          <source>Child Abuse Negl</source>
          <year>2023</year>
          <volume>143</volume>
          <fpage>106333</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/37379728"/>
          </comment>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2023.106333</pub-id>
          <pub-id pub-id-type="medline">37379728</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(23)00321-6</pub-id>
          <pub-id pub-id-type="pmcid">PMC10651183</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref56">
        <label>56</label>
        <nlm-citation citation-type="web">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Dale</surname>
              <given-names>M</given-names>
            </name>
          </person-group>
          <article-title>Addressing the underlying issue of poverty in child-neglect cases</article-title>
          <source>American Bar Association</source>
          <year>2014</year>
          <access-date>2025-03-19</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.americanbar.org/groups/litigation/resources/newsletters/childrens-rights/addressing-underlying-issue-poverty-child-neglect-cases/">https://www.americanbar.org/groups/litigation/resources/newsletters/childrens-rights/addressing-underlying-issue-poverty-child-neglect-cases/</ext-link>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref57">
        <label>57</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Rochford</surname>
              <given-names>HI</given-names>
            </name>
            <name name-style="western">
              <surname>Zeiger</surname>
              <given-names>KD</given-names>
            </name>
            <name name-style="western">
              <surname>Peek-Asa</surname>
              <given-names>C</given-names>
            </name>
          </person-group>
          <article-title>State-level education policies: opportunities for secondary prevention of child maltreatment</article-title>
          <source>Child Abuse Negl</source>
          <year>2023</year>
          <volume>136</volume>
          <fpage>106018</fpage>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2022.106018</pub-id>
          <pub-id pub-id-type="medline">36630852</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(22)00552-X</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref58">
        <label>58</label>
        <nlm-citation citation-type="book">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Shusterman</surname>
              <given-names>GR</given-names>
            </name>
            <name name-style="western">
              <surname>Hollinshead</surname>
              <given-names>D</given-names>
            </name>
            <name name-style="western">
              <surname>Fluke</surname>
              <given-names>JD</given-names>
            </name>
            <name name-style="western">
              <surname>Yuan</surname>
              <given-names>YT</given-names>
            </name>
          </person-group>
          <source>Alternative Responses to Child Maltreatment: Findings From NCANDS</source>
          <year>2005</year>
          <publisher-loc>USA</publisher-loc>
          <publisher-name>U.S. Department of Health and Human Services; Office of the Assistant Secretary for Planning and Evaluation</publisher-name>
        </nlm-citation>
      </ref>
      <ref id="ref59">
        <label>59</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Font</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Maguire-Jack</surname>
              <given-names>K</given-names>
            </name>
          </person-group>
          <article-title>The organizational context of substantiation in child protective services cases</article-title>
          <source>J Interpers Violence</source>
          <year>2021</year>
          <volume>36</volume>
          <issue>15-16</issue>
          <fpage>7414</fpage>
          <lpage>7435</lpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/30862238"/>
          </comment>
          <pub-id pub-id-type="doi">10.1177/0886260519834996</pub-id>
          <pub-id pub-id-type="medline">30862238</pub-id>
          <pub-id pub-id-type="pmcid">PMC7430033</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref60">
        <label>60</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Saar-Heiman</surname>
              <given-names>Y</given-names>
            </name>
          </person-group>
          <article-title>Understanding the relationships among poverty, child maltreatment, and child protection involvement: perspectives of service users and practitioners</article-title>
          <source>J Soc Social Work Res</source>
          <year>2022</year>
          <volume>13</volume>
          <issue>1</issue>
          <fpage>117</fpage>
          <lpage>141</lpage>
          <pub-id pub-id-type="doi">10.1086/713999</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref61">
        <label>61</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Lefebvre</surname>
              <given-names>R</given-names>
            </name>
            <name name-style="western">
              <surname>Fallon</surname>
              <given-names>B</given-names>
            </name>
            <name name-style="western">
              <surname>Van Wert</surname>
              <given-names>M</given-names>
            </name>
            <name name-style="western">
              <surname>Filippelli</surname>
              <given-names>J</given-names>
            </name>
          </person-group>
          <article-title>Examining the relationship between economic hardship and child maltreatment using data from the Ontario incidence study of reported child abuse and neglect-2013 (OIS-2013)</article-title>
          <source>Behav Sci (Basel)</source>
          <year>2017</year>
          <volume>7</volume>
          <issue>1</issue>
          <fpage>6</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.mdpi.com/resolver?pii=bs7010006"/>
          </comment>
          <pub-id pub-id-type="doi">10.3390/bs7010006</pub-id>
          <pub-id pub-id-type="medline">28208690</pub-id>
          <pub-id pub-id-type="pii">bs7010006</pub-id>
          <pub-id pub-id-type="pmcid">PMC5371750</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref62">
        <label>62</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Covington</surname>
              <given-names>T</given-names>
            </name>
            <name name-style="western">
              <surname>Collier</surname>
              <given-names>A</given-names>
            </name>
          </person-group>
          <article-title>Child maltreatment fatality reviews: learning together to improve systems that protect children and prevent maltreatment</article-title>
          <source>National Center for Fatality Review and Prevention</source>
          <year>2018</year>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.ncfrp.org/wp-content/uploads/NCRPCD-Docs/CAN_Guidance.pdf"/>
          </comment>
        </nlm-citation>
      </ref>
      <ref id="ref63">
        <label>63</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Campbell</surname>
              <given-names>KA</given-names>
            </name>
            <name name-style="western">
              <surname>Wood</surname>
              <given-names>JN</given-names>
            </name>
            <name name-style="western">
              <surname>Lindberg</surname>
              <given-names>DM</given-names>
            </name>
            <name name-style="western">
              <surname>Berger</surname>
              <given-names>RP</given-names>
            </name>
          </person-group>
          <article-title>A standardized definition of near-fatal child maltreatment: results of a multidisciplinary Delphi process</article-title>
          <source>Child Abuse Negl</source>
          <year>2021</year>
          <volume>112</volume>
          <fpage>104893</fpage>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://europepmc.org/abstract/MED/33373847"/>
          </comment>
          <pub-id pub-id-type="doi">10.1016/j.chiabu.2020.104893</pub-id>
          <pub-id pub-id-type="medline">33373847</pub-id>
          <pub-id pub-id-type="pii">S0145-2134(20)30548-2</pub-id>
          <pub-id pub-id-type="pmcid">PMC7856008</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref64">
        <label>64</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Rehill</surname>
              <given-names>P</given-names>
            </name>
            <name name-style="western">
              <surname>Biddle</surname>
              <given-names>N</given-names>
            </name>
          </person-group>
          <article-title>Transparency challenges in policy evaluation with causal machine learning: improving usability and accountability</article-title>
          <source>Data Policy</source>
          <year>2024</year>
          <volume>6</volume>
          <fpage>e43</fpage>
          <pub-id pub-id-type="doi">10.1017/dap.2024.35</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref65">
        <label>65</label>
        <nlm-citation citation-type="journal">
          <person-group person-group-type="author">
            <name name-style="western">
              <surname>Tutsoy</surname>
              <given-names>O</given-names>
            </name>
            <name name-style="western">
              <surname>Polat</surname>
              <given-names>A</given-names>
            </name>
            <name name-style="western">
              <surname>Colak</surname>
              <given-names>S</given-names>
            </name>
            <name name-style="western">
              <surname>Balikci</surname>
              <given-names>K</given-names>
            </name>
          </person-group>
          <article-title>Development of a multi-dimensional parametric model with non-pharmacological policies for predicting the COVID-19 pandemic casualties</article-title>
          <source>IEEE Access</source>
          <year>2020</year>
          <volume>8</volume>
          <fpage>225272</fpage>
          <lpage>225283</lpage>
          <pub-id pub-id-type="doi">10.1109/access.2020.3044929</pub-id>
        </nlm-citation>
      </ref>
      <ref id="ref66">
        <label>66</label>
        <nlm-citation citation-type="web">
          <source>National Data Archive on Child Abuse and Neglect</source>
          <access-date>2026-08-27</access-date>
          <comment>
            <ext-link ext-link-type="uri" xlink:type="simple" xlink:href="https://www.ndacan.acf.hhs.gov/">https://www.ndacan.acf.hhs.gov/</ext-link>
          </comment>
        </nlm-citation>
      </ref>
    </ref-list>
  </back>
</article>
