<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="research-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Ment Health</journal-id><journal-id journal-id-type="publisher-id">mental</journal-id><journal-id journal-id-type="index">16</journal-id><journal-title>JMIR Mental Health</journal-title><abbrev-journal-title>JMIR Ment Health</abbrev-journal-title><issn pub-type="epub">2368-7959</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v13i1e105460</article-id><article-id pub-id-type="doi">10.2196/105460</article-id><article-categories><subj-group subj-group-type="heading"><subject>Original Paper</subject></subj-group></article-categories><title-group><article-title>Assessing the Need for Mental Health Support From Free-Text Responses: Development and Validation of Language-Based Assessments in Adults With Internalizing Symptoms</article-title></title-group><contrib-group><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Wiebel</surname><given-names>Clara</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" corresp="yes" equal-contrib="yes"><name name-style="western"><surname>Eijsbroek</surname><given-names>Veerle C</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Varadarajan</surname><given-names>Vasudha</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff3">3</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Kjell</surname><given-names>Katarina</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Schwartz</surname><given-names>H Andrew</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff4">4</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Kjell</surname><given-names>Oscar</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff4">4</xref></contrib></contrib-group><aff id="aff1"><institution>Department of Psychology, Lund University</institution><addr-line>Box 213, Allhelgona Kyrkogata 16A, 16B, 18A, 18B och 18C</addr-line><addr-line>Lund</addr-line><addr-line>Sk&#x00E5;ne</addr-line><country>Sweden</country></aff><aff id="aff2"><institution>School of Management, Technical University of Munich</institution><addr-line>Heilbronn</addr-line><addr-line>Baden-W&#x00FC;rttemberg</addr-line><country>Germany</country></aff><aff id="aff3"><institution>Language Technologies Institute, Carnegie Mellon University</institution><addr-line>Pittsburgh</addr-line><addr-line>PA</addr-line><country>United States</country></aff><aff id="aff4"><institution>College of Connected Computing, Vanderbilt University</institution><addr-line>Nashville</addr-line><addr-line>TN</addr-line><country>United States</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Torous</surname><given-names>John</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Nook</surname><given-names>Erik C</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Mesquiti</surname><given-names>Steven</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Veerle C Eijsbroek, MSc, Department of Psychology, Lund University, Box 213, Allhelgona Kyrkogata 16A, 16B, 18A, 18B och 18C, Lund, Sk&#x00E5;ne, 22350, Sweden, +46 46 222 00 00; <email>veerle.eijsbroek@psy.lu.se</email></corresp><fn fn-type="equal" id="equal-contrib1"><label>*</label><p>these authors contributed equally</p></fn></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>25</day><month>9</month><year>2026</year></pub-date><volume>13</volume><elocation-id>e105460</elocation-id><history><date date-type="received"><day>24</day><month>06</month><year>2026</year></date><date date-type="rev-recd"><day>11</day><month>08</month><year>2026</year></date><date date-type="accepted"><day>12</day><month>08</month><year>2026</year></date></history><copyright-statement>&#x00A9; Clara Wiebel, Veerle C Eijsbroek, Vasudha Varadarajan, Katarina Kjell, H Andrew Schwartz, Oscar Kjell. Originally published in JMIR Mental Health (<ext-link ext-link-type="uri" xlink:href="https://mental.jmir.org">https://mental.jmir.org</ext-link>), 25.9.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Mental Health, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mental.jmir.org/">https://mental.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mental.jmir.org/2026/1/e105460"/><abstract><sec><title>Background</title><p>Machine learning and natural language processing have demonstrated significant potential for mental health assessment: describing your mental health in your own words can offer a more ecologically valid approach than traditional rating scales. However, most models focus on specific diagnoses, conditions, or symptoms, which may prematurely assign labels and potentially reinforce stigma in the context of early-stage mental health screening.</p></sec><sec><title>Objective</title><p>This study develops a language-based assessment model that assesses the need for mental health support based on probed natural language and validates it against best-estimate assessments from multiple experienced psychotherapists.</p></sec><sec sec-type="methods"><title>Methods</title><p>We analyzed an enriched online sample (n=600 for development and n=212 for validation), in which about half reported experiencing internalizing symptoms (depression or anxiety). Participants described their mental health using open-ended responses regarding (1) mental health, (2) suicidal thoughts, (3) medical history, and (4) depression. The responses were converted into contextual word embeddings using a large language model and entered as predictors in a ridge regression using nested cross-validation. Two to three experienced psychotherapists assessed each participant&#x2019;s need for mental health support on a scale from 1 (no support needed) to 5 (potential crisis). Their assessments were based on longitudinal clinical data (natural language, validated scales, clinical interview, sociodemographics, and clinical history) and were averaged into a best-estimate assessment for model validation. We used the Sequential Evaluation With Model Preregistration framework, which separates model development from validation in a held-out set to support robust estimations and generalizability.</p></sec><sec sec-type="results"><title>Results</title><p>The language-based assessments closely aligned with the best-estimate assessments (<italic>r</italic>=0.82) and showed strong correlations with established clinical rating scales for depression (Patient Health Questionnaire-9), anxiety (Generalized Anxiety Disorder 7-Item Scale), stress (Perceived Stress Scale 10), and suicidality (Inventory of Depression and Anxiety Symptoms; <italic>r</italic>=0.62-0.77). Language-based visualizations of topics and word embeddings showed that low need for support assessments was associated with mentioning well-being and good health, while high assessments were related to depression, anxiety, and suicidality.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>This study demonstrates that natural language responses analyzed through large language models and machine learning can be used to assess individuals&#x2019; need for mental health support in close alignment with best-estimate assessments from experienced psychotherapists. Using less than 5 minutes of respondent time, this approach offers a practical tool for early-stage mental health screening in both clinical and self-guided screening contexts.</p></sec></abstract><kwd-group><kwd>mental health support</kwd><kwd>language-based assessments</kwd><kwd>natural language processing</kwd><kwd>machine learning</kwd><kwd>large language models</kwd><kwd>best-estimate assessments</kwd><kwd>AI</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Nearly half of EU (European Union) residents reported emotional distress within the past year, such as feeling depressed or anxious [<xref ref-type="bibr" rid="ref1">1</xref>]. Yet, even in high-income countries, only 23% of people with mental health issues receive minimally adequate treatment [<xref ref-type="bibr" rid="ref2">2</xref>]. Personal barriers play a key role in this treatment gap: across 6 countries in North and South America, nearly 40% of people who did not seek any treatment for their mental health issues reported that they did not perceive a need for it [<xref ref-type="bibr" rid="ref3">3</xref>]. Among those who recognized a need, the most common barriers were attitudinal, such as believing that their issues were not serious enough.</p><p>These patterns reflect what the literature terms unmet need for mental health care, which can arise at 3 stages on the pathway to care: not perceiving a need, not seeking care, and not receiving adequate care [<xref ref-type="bibr" rid="ref4">4</xref>]. The first stage is pivotal, since perceiving a need strongly shapes whether people seek and use services [<xref ref-type="bibr" rid="ref5">5</xref>]. Yet, perceived need is captured through self-report instruments such as the Perceived Need for Care Questionnaire [<xref ref-type="bibr" rid="ref6">6</xref>] that depend on individuals recognizing a need themselves. This self-recognition is often absent precisely where clinical need is present: even among people who meet diagnostic criteria, those who do not perceive a need for care nonetheless show elevated distress and impairment [<xref ref-type="bibr" rid="ref7">7</xref>]. There is thus a gap for accessible, early-stage assessments of support needs that do not depend on the individual already recognizing that need.</p><p>Someone&#x2019;s need for mental health support can arise from a wide array of biological (eg, genes), developmental (eg, aging), psychological (eg, coping style), social (eg, loneliness), and structural (eg, discrimination) factors [<xref ref-type="bibr" rid="ref8">8</xref>-<xref ref-type="bibr" rid="ref10">10</xref>]. Rather than constituting need directly, these factors are determinants that shape the need for mental health support largely by contributing to psychological distress. In the current study, we assess the need for mental health support within the scope of internalizing symptoms (depression or anxiety) and suicidality. While internalizing symptoms do not encompass all possible manifestations of psychological distress (eg, externalizing symptoms), they represent one of the most prevalent, consistently used, and validated indicators of psychological distress in both research and clinical practice, while suicidality functions as a key marker of distress severity and urgency of need [<xref ref-type="bibr" rid="ref11">11</xref>-<xref ref-type="bibr" rid="ref13">13</xref>].</p><p>One promising approach for assessing need for mental health support is language-based assessments (LBAs), where an individual&#x2019;s language use is analyzed to assess their mental state. The analysis of language use patterns&#x2014;how people express their emotions, thoughts, or experiences in their own words&#x2014;has provided valuable psychological insights for over multiple decades [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref15">15</xref>]. Recent advances in natural language processing and machine learning, particularly large language models (LLMs), have substantially improved the accuracy of LBAs [<xref ref-type="bibr" rid="ref16">16</xref>,<xref ref-type="bibr" rid="ref17">17</xref>]. Accordingly, a growing body of work has applied natural language processing and LLMs to assess and understand mental illness. Within the scope of internalizing symptoms and suicidality, studies have analyzed social media posts [<xref ref-type="bibr" rid="ref18">18</xref>,<xref ref-type="bibr" rid="ref19">19</xref>], text responses and essays [<xref ref-type="bibr" rid="ref20">20</xref>,<xref ref-type="bibr" rid="ref21">21</xref>], and text messages and conversations [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref23">23</xref>], as well as clinical interview and psychotherapy transcripts [<xref ref-type="bibr" rid="ref24">24</xref>-<xref ref-type="bibr" rid="ref26">26</xref>], collectively yielding meaningful insight into depression, anxiety, and suicidality symptoms, profiles, and trajectories. A related line of research has explored probed LBAs, where individuals describe their symptoms and related situations in response to targeted open-ended mental health questions [<xref ref-type="bibr" rid="ref27">27</xref>]. Probed LBAs have been proposed as an ecologically valid complement to traditional assessment methods because they offer respondents more flexibility than closed-ended rating scales [<xref ref-type="bibr" rid="ref28">28</xref>,<xref ref-type="bibr" rid="ref29">29</xref>]. LLM-based analysis of probed responses has achieved convergent validity with rating scales approaching the scales&#x2019; own reliability&#x2014;a theoretical upper limit of assessment accuracy [<xref ref-type="bibr" rid="ref30">30</xref>,<xref ref-type="bibr" rid="ref31">31</xref>].</p><p>Despite these advances, current LBAs are restricted in several ways in relation to assessing the need for mental health support. First, most models focus on specific diagnoses, conditions, or symptoms, which may disregard psychiatric comorbidities and prematurely assign diagnostic labels, potentially reinforcing stigma and discouraging help-seeking [<xref ref-type="bibr" rid="ref32">32</xref>,<xref ref-type="bibr" rid="ref33">33</xref>]. More recently, a smaller set of studies has moved toward broader dimensional or transdiagnostic framings [<xref ref-type="bibr" rid="ref12">12</xref>], modeling spectra such as internalizing distress, rather than predicting diagnostic categories [<xref ref-type="bibr" rid="ref34">34</xref>,<xref ref-type="bibr" rid="ref35">35</xref>]. While these approaches can better capture comorbidity and reduce reliance on diagnostic labels, they remain oriented toward characterizing symptom dimensions or disorder spectra rather than directly assessing an individual&#x2019;s need for support. While assessing internalizing symptoms and especially suicidality directly relates to the need for mental health support, we argue that assessing overall need first, rather than focusing on a specific condition or spectrum, can avoid these issues. It is potentially better suited for early-stage screening, where a formal diagnosis or condition-specific severity estimation is neither necessary nor appropriate. Although LBAs have successfully advanced the identification of internalizing symptoms and suicidality, to our knowledge, no validated instrument to date screens for the need for mental health support using natural language without relying on a specific diagnosis, condition, or spectrum. Meanwhile, broader screening tools such as the General Health Questionnaire [<xref ref-type="bibr" rid="ref36">36</xref>] rely exclusively on rating scales.</p><p>Second, many LBAs are trained on passively collected language, such as social media posts [<xref ref-type="bibr" rid="ref18">18</xref>], text messages [<xref ref-type="bibr" rid="ref22">22</xref>], medical records [<xref ref-type="bibr" rid="ref37">37</xref>], or interview or therapy transcripts [<xref ref-type="bibr" rid="ref25">25</xref>]. While valuable, these sources are not universally accessible or applicable for early screening: not everyone maintains a social media profile or produces sufficient digital text for analysis; personal messages raise privacy concerns; medical records vary in availability and format across health care systems; and interview or therapy transcripts are only generated once someone has already moved past the point of early screening. Probed language responses, where individuals answer the same targeted questions, provide more standardized input that isolates the constructs of interest (and typically align the language source in time and content with the outcome, producing high convergent validity [<xref ref-type="bibr" rid="ref28">28</xref>,<xref ref-type="bibr" rid="ref30">30</xref>]).</p><p>Third, models are typically validated against a reference standard based on a single measure, such as rating scales [<xref ref-type="bibr" rid="ref38">38</xref>] or structured interviews [<xref ref-type="bibr" rid="ref39">39</xref>]. This limits their validity, since, in essence, every single measure of a psychological construct is prone to some sort of error or bias, such as recall bias in self-report scales [<xref ref-type="bibr" rid="ref40">40</xref>] or interviewer bias in clinical interviews [<xref ref-type="bibr" rid="ref41">41</xref>], and therefore they are unable to be a best estimate of a clinical construct [<xref ref-type="bibr" rid="ref42">42</xref>,<xref ref-type="bibr" rid="ref43">43</xref>].</p><p>Finally, even when machine learning models perform well on training data, they often fail to generalize to unseen data, which is a problem well-documented across predictive modeling in clinical science [<xref ref-type="bibr" rid="ref44">44</xref>], suggesting that without rigorous validation procedures, such as testing preregistered models on held-out samples, reported accuracies may be overly optimistic.</p><p>This study addresses several limitations of previous approaches. First, it introduces assessments of the need for mental health support as an early complement or alternative to single-label diagnoses or severity estimations of specific conditions or symptoms. Need for support aims to account for comorbid mental illnesses and offers a less prescriptive approach. Conceptually, the assessments resemble the judgment of a mental health professional right after an initial consultation&#x2014;on a clinician-rated scale ranging from (1) no support needed to (5) potential crisis. As the output is framed as supportive guidance rather than a diagnostic label or condition-specific severity estimation, it has the potential to be suited not only for clinical screening and monitoring, but also for guided self-screening contexts such as digital health platforms, where individuals can receive advice framed as supportive rather than clinical, indicating whether seeking professional support may be beneficial.</p><p>Second, our approach relies on structured, probed language responses rather than passively collected language. Respondents answer targeted mental health questions, either by writing brief paragraphs or selecting descriptive words from a predefined list [<xref ref-type="bibr" rid="ref45">45</xref>]. This structured input produces standardized and comparable data, enhancing real-world applicability. It also enables examining whether a single general mental health question is sufficient to capture most needs, or whether additional questions, such as explicitly asking about medical history or suicidality, are needed to improve accuracy.</p><p>Third, the LBAs are validated against best-estimate assessments instead of a single measure. These best-estimates are obtained through a Longitudinal Expert All Data (LEAD) approach [<xref ref-type="bibr" rid="ref46">46</xref>,<xref ref-type="bibr" rid="ref47">47</xref>]: multiple experts assessed a case based on longitudinal clinical data (natural language, validated rating scales, clinical interview, sociodemographics, and clinical history), and their assessments are averaged to a &#x201C;best-estimate&#x201D; of the need for mental health support. Finally, to produce a fair estimation of replicability and generalizability, the current study adheres to the Sequential Evaluation With Model Preregistration (SEMP) framework [<xref ref-type="bibr" rid="ref48">48</xref>]. In the development phase, the optimal combination of input variables that balances accuracy and parsimony is identified and preregistered. In the validation phase, preregistered models and hypotheses are tested in a separate held-out set, evaluating the LBAs model&#x2019;s criterion, convergent, and face validity.</p><p>We hypothesize that the LBAs show strong criterion validity by correlating positively with best-estimate assessments from experienced psychotherapists as reference standard (H1). For convergent validity, we expect the LBAs to correlate positively with established measures of depression (H2a), anxiety (H2b), stress (H2c), and suicidality (H2d), and negatively with satisfaction with life (H2e) and harmony in life (H2f). For external validity, we expect positive correlations with self-reported sick days (H3a) and health care visits due to mental health (H3b) as behavioral indicators of mental health service use. Finally, we expect the LBAs to show face validity, with linguistic topics aligning with established psychological theories, specifically the broaden-and-build theory [<xref ref-type="bibr" rid="ref49">49</xref>] regarding low need for support and the theory of depression [<xref ref-type="bibr" rid="ref50">50</xref>,<xref ref-type="bibr" rid="ref51">51</xref>] regarding high need for support (H4). The broaden-and-build theory [<xref ref-type="bibr" rid="ref49">49</xref>] suggests that positive emotions expand cognitive and behavioral repertoires, facilitating well-being and social connectedness. Therefore, we expect LBAs of a low need for support to be associated with topics referring to, for example, social relationships, meaningful activities, recovery from stressors, a lack of mental or physical issues, and explicit affirmations of well-being. Conversely, the theory of depression by Beck [<xref ref-type="bibr" rid="ref50">50</xref>] emphasizes that negative thought patterns and attentional biases toward distress lead to negative affect, which Clark and Watson [<xref ref-type="bibr" rid="ref51">51</xref>] identify as a key feature of both depression and anxiety. Accordingly, we expect LBAs of a high need for support to be associated with topics reflecting negative thoughts, negative affect, and distressing life circumstances.</p></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Procedure</title><p>The participants&#x2019; mental health data used to develop and validate the LBA model were drawn from a previously conducted longitudinal study focusing on assessing depression, anxiety, stress, suicidality, and self-harm [<xref ref-type="bibr" rid="ref45">45</xref>]. Participants completed a comprehensive assessment survey at the beginning and end of the current study&#x2019;s period, as well as shorter biweekly surveys in between. The surveys assessed mental and physical health through 3 response formats: brief written paragraphs (text responses), writing or selecting words from a list (word responses), and standardized rating scales. The order of the response formats (ie, open-ended questions or rating scales first) was randomized. The median estimated response time for the 4 language responses included in the final model (see the Model Development section) was approximately 4.79 (mean 6.14, SD 7.84) minutes.</p></sec><sec id="s2-2"><title>Open Science</title><p>In accordance with the SEMP framework, the first preregistration [<xref ref-type="bibr" rid="ref52">52</xref>] outlines the model development, student assessment procedure, sample size, and exclusion criteria. After model development, a second preregistration [<xref ref-type="bibr" rid="ref53">53</xref>] specified the final models and refined hypotheses tested in a held-out set. The final models are available in the LBA model library [<xref ref-type="bibr" rid="ref54">54</xref>] and can be explored at our public demo page [<xref ref-type="bibr" rid="ref55">55</xref>].</p></sec><sec id="s2-3"><title>Participants</title><p>Participants from the United States and United Kingdom were recruited via Prolific [<xref ref-type="bibr" rid="ref56">56</xref>] and followed online for 10 weeks between June and December 2021. The participant sample was enriched, with approximately half of the participants self-reporting having been diagnosed with a mental health diagnosis of major depressive disorder (MDD [<xref ref-type="bibr" rid="ref57">57</xref>]; &#x2248;25%) or Generalized Anxiety Disorder (GAD [<xref ref-type="bibr" rid="ref57">57</xref>]; &#x2248;25%). Of the total sample, 212 participants were assigned to a held-out validation set based on having received assessments from multiple psychotherapists, and 600 were assigned to the development set (see the Selection of Development Sample section).</p><p>The sample was predominantly female (development: 396/600=66%; validation: 114/212=53.8%) with a mean age of 43.9 (SD 18.2) years in the development set and 40.2 (SD 13.0) years in the validation set. Participants in the validation set were, on average, slightly younger, more likely to hold a university degree, more likely employed, and more often male compared to the development set. Full sociodemographic details are presented in <xref ref-type="fig" rid="figure1">Figure 1</xref> and <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>Sociodemographics in the development (n=600) and validation set (n=212). Participants could choose multiple options for the occupation variable.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e105460_fig01.png"/></fig></sec><sec id="s2-4"><title>Measures</title><sec id="s2-4-1"><title>Natural Language Responses</title><p>Six text response questions and 2 word response questions were chosen from Kjell et al [<xref ref-type="bibr" rid="ref45">45</xref>] based on clinical expertise being most relevant for determining need for mental health support and used as possible inputs for the LBA model. The six text response questions prompted participants to write brief paragraphs about their (1) mental health, (2) social situation, (3) support system, (4) medical history, (5) self-harm, and (6) suicidality. The two questions for word responses prompted participants to choose at least five words from a list of 20 words that describe their (1) mental health and (2) depressive symptoms over the past two weeks. Additionally, participants could also write their own descriptive words in an open text field. The exact wording of the questions included in the final model is presented below (see Table S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> for all text and word response questions).</p><p>The mental health text response question asked:</p><disp-quote><p>How is your mental health? Please describe how you have been over the last two weeks. You can, for example, write about your emotions, thoughts, behaviors, and/or symptoms related to your health. Write at least one paragraph.</p></disp-quote><p>The suicidality text response question asked:</p><disp-quote><p>Please describe whether, and if so how, you have been thinking about death, have had thoughts about killing yourself, have any intent or plan to kill yourself, or if you have attempted suicide over the last two weeks? If you have not had or done this, please briefly describe that in your own words.</p></disp-quote><p>The medical history text response question asked:</p><disp-quote><p>Describe your previous medical history; for example, when did you get sick, have you had similar or other problems before?</p></disp-quote><p>The depression word response question asked:</p><disp-quote><p>Write or select five words that best describe the level of depression that you have experienced, or not, over the last two weeks.</p><p>Please choose all that apply: content, motivated, optimistic, happy, hopeful, fearful, energetic, lazy, okay, upbeat, suicidal, anxiety, angry, alone, numb, pessimistic, empty, lethargic, dejected, despondent, other: (specify).</p></disp-quote></sec><sec id="s2-4-2"><title>Rating Scales and Behavioral Measures</title><p>The Patient Health Questionnaire-9 (PHQ-9) [<xref ref-type="bibr" rid="ref58">58</xref>] is a 9-item instrument that measures the severity of depression. The items correspond to the diagnostic criteria for MDD [<xref ref-type="bibr" rid="ref57">57</xref>]. Respondents rate how frequently they have experienced each symptom over the past 2 weeks on a scale of 0 (&#x201C;not at all&#x201D;) to 3 (&#x201C;nearly every day&#x201D;). An example item is &#x201C;Over the last two weeks, how often have you been bothered by little interest or pleasure in doing things?&#x201D;</p><p>The Generalized Anxiety Disorder 7-Item Scale (GAD-7) [<xref ref-type="bibr" rid="ref59">59</xref>] measures the severity of GAD symptoms on a 7-item scale. Respondents indicate how often they have experienced symptoms of GAD over the past 2 weeks on a scale of 0 (&#x201C;not at all&#x201D;) to 3 (&#x201C;nearly every day&#x201D;). An example item is &#x201C;Over the last two weeks, how often have you been bothered by feeling nervous, anxious, or on edge?&#x201D;</p><p>The Perceived Stress Scale 10 (PSS-10) [<xref ref-type="bibr" rid="ref60">60</xref>] is a 10-item instrument that assesses the frequency of subjectively experienced stress over the past month. Each item is answered on a scale of 0 (&#x201C;never&#x201D;) to 4 (&#x201C;very often&#x201D;), for instance, &#x201C;In the last month, how often have you felt difficulties were piling up so high that you could not overcome them?&#x201D;</p><p>The Inventory of Depression and Anxiety Symptoms (IDAS) [<xref ref-type="bibr" rid="ref61">61</xref>] assesses 10 symptom dimensions of MDD and related anxiety disorders using 64 items. Participants rate the extent to which they experienced each symptom in the past 2 weeks on a scale of 1 (&#x201C;not at all&#x201D;) to 5 (&#x201C;extremely&#x201D;). This study exclusively uses the suicidality dimension, which consists of 6 items; for example, &#x201C;I thought the world would be better off without me.&#x201D;</p><p>The Satisfaction With Life Scale 3 (SWLS-3) and the Harmony in Life Scale 3 (HILS-3) [<xref ref-type="bibr" rid="ref62">62</xref>] are abbreviated 3-item versions of their respective full scales, measuring global life satisfaction and harmony in life [<xref ref-type="bibr" rid="ref63">63</xref>,<xref ref-type="bibr" rid="ref64">64</xref>]. Both scales use a 7-point response format ranging from 1 (&#x201C;strongly disagree&#x201D;) to 7 (&#x201C;strongly agree&#x201D;). Example items include &#x201C;In most ways, my life is close to my ideal&#x201D; for the SWLS-3 and &#x201C;Most aspects of my life are in balance&#x201D; for the HILS-3.</p><p>The self-reported behavioral measures used in this study are the self-reported number of sick days and health care visits taken due to mental health issues over the past 3 months.</p></sec><sec id="s2-4-3"><title>Need for Mental Health Support</title><p>The Need for Mental Health Support Scale was developed to assess the need for support based on multiple sources of longitudinal data [<xref ref-type="bibr" rid="ref45">45</xref>]. It is completed by trained assessors (not self-reported)&#x2014;in this study, psychotherapists and graduate students of psychology. The assessors could choose between 5 levels ranging from 1 (no support needed) to 5 (potential crisis; <xref ref-type="table" rid="table1">Table 1</xref>). The tool was designed with 2 audiences in mind: psychotherapists assessing the level of support needed, and those assessed who may receive the output. Notably, item formulations 4 and 5 start similarly to avoid alarming language, as these might be shown to those assessed; the distinction between the 2 levels is conveyed through the accompanying descriptions, which include more urgency and emergency information for level 5.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Need for mental health support scale. The instruction to the assessors was &#x201C;Which recommendation would you give the patient based on their overall responses?&#x201D; The item formulations are addressed to those assessed because the resulting recommendation may later be shown to them.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Label</td><td align="left" valign="bottom">Item formulation</td></tr></thead><tbody><tr><td align="left" valign="top">No support needed</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>1=Your mental health appears to be well.</p></list-item><list-item><p>Your responses suggest that you do not need to consider seeking help for common mental health issues (ie, depression, anxiety, or stress).</p></list-item></list></td></tr><tr><td align="left" valign="top">Heightened attention</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>2=Consider paying attention to your mental health.</p></list-item><list-item><p>Your responses suggest that although you do not have to seek help for common mental health issues (ie, depression, anxiety, or stress) at this point, there are things you can do on your own to improve your mental health. You can, for example, exercise, meditate, eat well, be more mindful, etc. If things are getting worse, consider reaching out to healthcare.</p></list-item></list></td></tr><tr><td align="left" valign="top">Consider seeking support</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>3=Consider talking to a mental health clinician.</p></list-item><list-item><p>Talking to a mental health expert can help you increase your mental health before it gets worse. They can, for example, help you get the right help and give support.</p></list-item></list></td></tr><tr><td align="left" valign="top">Seek help (noncrisis)</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>4=Seek help from a mental health clinician.</p></list-item><list-item><p>Talking to mental health clinicians can help you get the right support; this may, for example, include therapy and/or medication, which research has shown can help decrease common mental health problems.</p></list-item><list-item><p>You may also get support from available helplines such as the Samaritans [<ext-link ext-link-type="uri" xlink:href="https://www.samaritans.org/">65</ext-link>], 116 123 in the UK or the National Suicide Prevention Lifeline [<xref ref-type="bibr" rid="ref66">66</xref>], 1-800-273-TALK (8255) in the US.</p></list-item></list></td></tr><tr><td align="left" valign="top">Potential crisis</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>5=Seek help from a mental health clinician.</p></list-item><list-item><p>Getting help from mental health clinicians can help you in many ways. If you need to reach immediate care, please contact your local emergency number such as 999 or 112.</p></list-item><list-item><p>You may also get support from available helplines such as the Samaritans [<xref ref-type="bibr" rid="ref65">65</xref>], 116 123 in the UK or the National Suicide Prevention Lifeline [<xref ref-type="bibr" rid="ref66">66</xref>], 1-800-273-TALK (8255) in the US.</p></list-item></list></td></tr></tbody></table></table-wrap></sec></sec><sec id="s2-5"><title>Assessment Procedure</title><sec id="s2-5-1"><title>Best-Estimate Assessments</title><p>Participants in the held-out validation set received a best-estimate assessment according to the LEAD method (described according to the LEADING reporting guideline [<xref ref-type="bibr" rid="ref47">47</xref>] in Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), which we used for model validation. As posed by Spitzer [<xref ref-type="bibr" rid="ref46">46</xref>], the LEAD method establishes a near-ground truth (&#x201C;best-estimate&#x201D; [<xref ref-type="bibr" rid="ref47">47</xref>]) for evaluating the validity of psychiatric assessments: participants are followed longitudinally and assessed by experts who have access to all collected data. Accordingly, experienced psychotherapists provided mental health assessments based on comprehensive, structured reports that included all information from the longitudinal surveys (eg, language responses, rating scales, sociodemographics, unstructured clinical interview, and clinical history; <xref ref-type="fig" rid="figure2">Figure 2</xref> [<xref ref-type="bibr" rid="ref28">28</xref>]). Each case was assessed by 2 or 3 licensed psychotherapists (&#x2265;5 y clinical experience; 2 women, 1 man; mean age 50.3, SD 12.9 y), and their assessments were averaged to form the best-estimate used as the reference standard.</p><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Overview of the student and best-estimate assessments for need of mental health support. For 167 cases in the development set, the student assessments were combined with a single psychotherapist assessment. GAD-7: Generalized Anxiety Disorder 7-Item Scale; PHQ-9: Patient Health Questionnaire 9-Item Scale.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e105460_fig02.png"/></fig></sec><sec id="s2-5-2"><title>Student Assessments</title><p>Participant responses in the development set (n=600) were assessed by graduate students of psychology (<xref ref-type="fig" rid="figure2">Figure 2</xref>). Each participant received 2 to 4 independent assessments, which were averaged for model training. Unlike the psychotherapists who reviewed all longitudinal data, the students based their assessments solely on the 6 text responses as well as the participant&#x2019;s age, gender, and occupation status. Moreover, they only saw responses from the first assessment survey, not longitudinal data.</p><p>The student assessments were collected via an online survey on SoSci Survey [<xref ref-type="bibr" rid="ref67">67</xref>] (Figure S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Each student assessed 40 participants, who were allocated semirandomly by the survey software so that each participant would receive an equal number of assessments. Four assessors (160 assessments) were excluded for completing the survey in under 15 (median completion time 45.03, IQR 35.57-80.60) minutes. The number of assessments per participant therefore varied between 2 and 4 (in the development set: 188, 31.3%, received two assessments; 392, 65.3%, three assessments; and 20, 3.3%, received four assessments).</p><p>The 51 final graduate student assessors were aged between 21 and 32 (mean 25.63, SD 2.62) years; 39 (76.5%) students were female, 11 (21.6%) students were male, and 1 student identified as other. A substantial majority (n=42 or 82.4%) had previous experience with patients with psychiatric disorders, either in an inpatient (n=29 or 56.9%) or outpatient (n=25 or 49%) setting. Moreover, 36 (70.6%) assessors reported previous experience with psychological diagnostics, either in clinical (n=31 or 60.8%) or nonclinical settings (n=8 or 15.7%).</p><p>To evaluate how similar the best-estimate and student assessments were, the validation set (n=212) was also assessed by the graduate students. We calculated three interrater reliability metrics in the validation set using Krippendorff &#x03B1;: reliability (1) among the psychotherapists, (2) among the students, and (3) between the best-estimate and averaged student assessments.</p></sec></sec><sec id="s2-6"><title>Ethical Considerations</title><p>This study received ethical approval from the Swedish National Ethics Committee (Dnr 2021&#x2010;01820). Participants were informed about this study and that participation was voluntary. Consent was obtained before starting this study. Where participants answered open-ended questions on suicide and self-harm, they were required to tick a box indicating that they understood that the survey was anonymous and that no one would be able to contact them in relation to their responses, and they were provided information about where they could retrieve help (eg, suicide prevention helplines). They were compensated through money (&#x00A3;7.5/h; a currency exchange rate of GB &#x00A3;1=US $1.37 was applicable) on Prolific. The data collection of student assessments was approved by the Ethics Committee of the University of Mainz (2024-JGU-psychEK-022). Student assessors were compensated through course credit.</p></sec><sec id="s2-7"><title>Model Development</title><sec id="s2-7-1"><title>Selection of Development Sample</title><p>The development set was drawn from the 1095 participants not included in the held-out validation set. Of those, 167 participants who had received one psychotherapist&#x2019;s assessment were included, and the remaining 433 were randomly selected, yielding a total of 600. This sample size was chosen to balance training accuracy with feasible annotation effort, based on Gu et al [<xref ref-type="bibr" rid="ref29">29</xref>], who showed that LBA models trained with 500 achieved accuracy within <italic>r</italic> =&#x00B1;0.05 of models trained with 963.</p></sec><sec id="s2-7-2"><title>Model Configuration</title><p>Selected language responses (see the Selection of Language Variables section) were transformed into word embeddings and used as predictors in a ridge regression to assess the need for mental health support. Specifically, the second-to-last layer of the transformer-based LLM mxbai-embed-large-v1 [<xref ref-type="bibr" rid="ref68">68</xref>] was used to generate the embeddings (<xref ref-type="fig" rid="figure3">Figure 3</xref>). This is a document-tuned transformer model, a type that has demonstrated state-of-the-art performance on the Massive Text Embedding Benchmark [<xref ref-type="bibr" rid="ref69">69</xref>] and has been shown to consistently outperform base transformer representations for person-level psychological assessment [<xref ref-type="bibr" rid="ref70">70</xref>,<xref ref-type="bibr" rid="ref71">71</xref>]. The model (Mixedbread AI [<xref ref-type="bibr" rid="ref68">68</xref>]; &#x2248;335 million parameters; 1024-dimensional embeddings) was used off the shelf without fine-tuning and run locally via the <italic>text</italic> package [<xref ref-type="bibr" rid="ref72">72</xref>]. A completed GUIDE-LLM checklist for reporting LLM use [<xref ref-type="bibr" rid="ref73">73</xref>] is provided in the supplementary material (Table S3 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>The embedding process for representing a participant&#x2019;s language responses. Language responses are processed using the mxbai-embed-large-v1 model. Each response is tokenized, and a 1024-dimensional embedding is generated for each token. Token embeddings are then averaged across tokens to form a single representation for each language response. Finally, embeddings from different language responses are concatenated to create a unified representation for all responses. LLM: large language model.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e105460_fig03.png"/></fig><p>The ridge regression was trained using 10-fold nested cross-validation to optimize the regularization parameter that minimizes overfitting. The search range was set from 10<sup>-6</sup> to 10<sup>6</sup>. The model&#x2019;s predictive accuracy was assessed using the Pearson correlation between actual and predicted values in the outer test folds. For some cases, the final LBAs were slightly below 1 or above 5, which were truncated to 1 and 5 for subsequent analyses.</p><p>During model development, we explored 3 methods to optimize word embeddings and improve model accuracy: averaging embeddings, Matryoshka embeddings [<xref ref-type="bibr" rid="ref74">74</xref>], and prepending the question to participant responses before embedding extraction.</p></sec><sec id="s2-7-3"><title>Selection of Language Variables</title><p>The selection of input variables followed an exploratory stepwise-forward approach. A ridge regression was fitted using only the embeddings of the mental health text responses since this question allowed respondents to broadly discuss diverse aspects of their mental health. Subsequently, each remaining variable was added individually to create multiple 2-variable LBA models. The LBA models were then repeatedly evaluated using data-driven and design-driven criteria until no further variables were selected. The primary criterion was data-driven: a variable was only considered if adding it significantly reduced model residuals, assessed via a 1-sided paired <italic>t</italic> test with &#x03B1;=.05. Although our first preregistration specified a Fisher z test as the data-driven criterion, we ultimately used the paired <italic>t</italic> test on model residuals due to its greater statistical power. Among variables reaching significance, design-driven considerations guided the final selection, including model parsimony, similarity to previously selected variables, and potential impact on the respondent&#x2019;s time burden.</p></sec></sec><sec id="s2-8"><title>Model Validation</title><p>Criterion validity (H1) was evaluated using the Pearson correlation between the LBAs and the best-estimate assessments in the validation set (n=212). A correlation of <italic>r</italic>&#x2265;0.70 was set as the threshold for adequate criterion validity [<xref ref-type="bibr" rid="ref75">75</xref>]. We additionally computed a 1-sided disattenuated correlation [<xref ref-type="bibr" rid="ref76">76</xref>]. Disattenuation estimates what the criterion validity correlation would be if the reference standard were measured without error&#x2014;that is, if the psychotherapists agreed perfectly on the need for mental health support. The correction is 1-sided because only the reference standard, not the LBA, is corrected for unreliability. We used the interrater reliability among the psychotherapists (Krippendorff <inline-formula><mml:math id="ieqn1"><mml:mi>&#x03B1;</mml:mi></mml:math></inline-formula>) as the correction factor. As error-free measurement is unattainable in practice, the disattenuated correlation should be interpreted as an upper-bound estimate rather than a point estimate.</p><p>Convergent validity (H2) was evaluated using Pearson correlations between the LBAs and rating scales in the validation set (n=212). We expected significant, at least medium-sized, correlations of |<italic>r</italic>|&#x003E;=0.30 [<xref ref-type="bibr" rid="ref77">77</xref>]; specifically, positive correlations with depression (PHQ-9; H2a), anxiety (GAD-7; H2b), stress (PSS-10; H2c), and suicidality (IDAS; H2d); and negative correlations with life satisfaction (SWLS-3; H2e) and harmony in life (HILS-3; H2f). External validity (H3) was evaluated using Pearson correlations between the LBAs and the behavioral measures. We expected significant, at least small (<italic>r</italic>&#x2265;0.10) [<xref ref-type="bibr" rid="ref77">77</xref>], positive correlations with the number of sick days (H3a) and the number of health care visits (H3b).</p><p>We also explored a parsimonious model using only the mental health text as a predictor and evaluated its criterion, convergent, and external validity (H1-H3). Such a model offers practical advantages, as open questions about mental health are commonly included in studies, interviews, and existing datasets, and requires less response time. The significance of all validity correlations (H1-H3) was 2-sided tested. Despite the directional hypotheses, we retained 2-sided tests as the more conservative choice.</p><p>Finally, face validity (H4) was assessed by examining how linguistic topics correlate with the LBAs. Unlike H1-H3, which evaluated model performance on held-out data, this analysis aimed to provide insights into the model&#x2019;s decisions through language patterns. Topics were extracted by applying latent Dirichlet allocation (LDA [<xref ref-type="bibr" rid="ref78">78</xref>]; see the Language-Based Visualizations section) on the complete dataset (n=812). LDA is an unsupervised method that is fully independent of the word embeddings, the ridge regression, and the human assessments. This justifies using the full sample, in addition to the fact that a larger sample yields more stable and interpretable topics. Yet, this analysis should be interpreted as descriptive face validity evidence rather than a test on unseen data. An a priori power analysis determined that with 812 participants, &#x03B1;=.05, and 90% power, correlations of |<italic>r</italic>|&#x2265;0.113 would be statistically significant [<xref ref-type="bibr" rid="ref79">79</xref>]. As topic-outcome correlations typically remain below |<italic>r</italic>|=0.30, we expected at least 4 significant topics per text variable and 2 significant topics per word variable.</p></sec><sec id="s2-9"><title>Language-Based Visualizations</title><p>We examined linguistic patterns associated with the LBAs using 2 complementary methods [<xref ref-type="bibr" rid="ref80">80</xref>]: LDA [<xref ref-type="bibr" rid="ref78">78</xref>] and supervised embedding projections [<xref ref-type="bibr" rid="ref72">72</xref>]. LDA is a topic modeling method that identifies latent topics within a text corpus based on word frequencies and co-occurrences. We extracted 12 topics per text variable and 2 topics per word variable. Before topic extraction, stop words [<xref ref-type="bibr" rid="ref81">81</xref>], repeated phrases from the questions, and highly frequent terms were removed, with frequency thresholds set individually per variable (ranging from 150 to 450 occurrences). Co-occurring words were structured into a Document Term Matrix, allowing for 1- to 3-grams. Topic distributions within each response were then used to predict the LBAs via linear regression, with standardized regression weights serving as effect sizes. False Discovery Rate corrections [<xref ref-type="bibr" rid="ref82">82</xref>] were applied within each language variable.</p><p>Supervised embedding projections [<xref ref-type="bibr" rid="ref72">72</xref>] provided a second, complementary visualization based on word embeddings rather than word frequency. This method identifies words significantly associated with the LBAs by comparing responses from the highest and lowest quartiles of the assessed need for support. Embeddings from both groups were aggregated, and a difference vector was computed. Each word was projected onto this vector using the dot product, mapping its association with the LBAs. Statistical significance was assessed against a permuted null distribution, with False Discovery Rate corrections [<xref ref-type="bibr" rid="ref82">82</xref>] applied within each variable. Stop words [<xref ref-type="bibr" rid="ref81">81</xref>] were removed before analysis.</p></sec><sec id="s2-10"><title>Statistical Software</title><p>All analyses were performed in R (version 4.4.1) [<xref ref-type="bibr" rid="ref83">83</xref>]. The model development and supervised embedding projections were performed using the <italic>text</italic> package (version 1.2.3) [<xref ref-type="bibr" rid="ref72">72</xref>]. LDA topics were extracted with the <italic>topics</italic> package (version 0.40.1) [<xref ref-type="bibr" rid="ref84">84</xref>]. Data wrangling, descriptive analyses, and visualizations were carried out with the packages <italic>dplyr</italic> (version 1.1.4) [<xref ref-type="bibr" rid="ref85">85</xref>] and <italic>ggplot</italic> (version 3.5.1) [<xref ref-type="bibr" rid="ref86">86</xref>]. Krippendorff &#x03B1; was calculated using the <italic>irr</italic> package (version 0.84.1) [<xref ref-type="bibr" rid="ref87">87</xref>].</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Descriptive Statistics</title><p>Participants wrote between a median of 9 with IQR 5-15 (self-harm) and 44 with IQR 30-68 words (mental health) in the text response variables (Table S4 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). A substantial proportion of participants (development: 240/600 or 40%; validation: 94/212 or 44.3%) fell within clinically relevant ranges of at least moderate depression and/or anxiety according to PHQ-9 and GAD-7 thresholds [<xref ref-type="bibr" rid="ref58">58</xref>,<xref ref-type="bibr" rid="ref59">59</xref>]. All rating scales demonstrated good to excellent internal consistency in both the development and validation sets (Cronbach &#x03B1;=0.85-0.96; see Table S5 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> for scale-specific values).</p><p>On average, participants&#x2019; need for mental health support in the validation set was 2.55 (SD 1.07) as assessed by the students and 2.50 (SD 1.16) as assessed by the psychotherapists; this difference was not significant, <italic>t</italic><sub>211</sub>=1.064, <italic>P</italic>=.29. An exploratory analysis of where the student and best-estimate assessments converged and diverged is reported in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>, as well as full descriptive statistics of the language responses, rating scales, and need for mental health support assessments (Tables S4 to S7 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p></sec><sec id="s3-2"><title>Model Development</title><p>Comparing embedding optimization methods revealed that prepending questions to responses before embedding extraction significantly improved model performance (see Table S8 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). This finding is theoretically plausible because contextual embeddings represent each token conditional on its surrounding text. Short responses in particular (especially in the suicidality text responses, eg, &#x201C;no, never&#x201D;) are often ambiguous in isolation, and prepending the corresponding question anchors the response to the construct being probed (eg, denial of suicidal ideation). Accordingly, this approach was adopted in all subsequent analyses.</p><p>Using only the mental health text response, the model achieved a training accuracy of <italic>r</italic>=0.79 (<xref ref-type="table" rid="table2">Table 2</xref>). In the first step of the step-forward procedure, adding suicidality yielded the largest increase in accuracy (<italic>r</italic>=0.79 to <italic>r</italic>=0.84). In the second step, medical history further improved accuracy (<italic>r</italic>=0.86). In the third step, only depression words significantly improved the model. Despite a small increase (&#x0394;<italic>r</italic>=0.01), we included it because participants typically select these words in under a minute. The final 4-variable model achieved a training accuracy of <italic>r</italic>=0.87.</p><table-wrap id="t2" position="float"><label>Table 2.</label><caption><p>Comparison of the selected models in the Stepwise-Forward approach (n=600). This table compares each model to the model in the row directly above through a paired, 1-sided <italic>t</italic> test of model residuals. See Table S9 (<xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) for all tested models.</p></caption><table id="table2" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Input variables</td><td align="left" valign="bottom"><italic>r</italic><sup><xref ref-type="table-fn" rid="table2fn1">a</xref></sup></td><td align="left" valign="bottom">95% CI</td><td align="left" valign="bottom">Mean abs residual<sup><xref ref-type="table-fn" rid="table2fn2">b</xref></sup> (SD)</td><td align="left" valign="bottom"><inline-formula><mml:math id="ieqn2"><mml:mo>&#x2206;</mml:mo></mml:math></inline-formula> res<sup><xref ref-type="table-fn" rid="table2fn3">c</xref></sup></td><td align="left" valign="bottom"><italic>t</italic> test (<italic>df</italic>)</td><td align="left" valign="bottom"><italic>P</italic> value</td></tr></thead><tbody><tr><td align="left" valign="top">Mental health</td><td align="left" valign="top">0.79</td><td align="left" valign="top">0.76-0.82</td><td align="left" valign="top">0.55 (0.42)</td><td align="left" valign="top">N/A<sup><xref ref-type="table-fn" rid="table2fn4">d</xref></sup></td><td align="left" valign="top">N/A</td><td align="left" valign="top">N/A</td></tr><tr><td align="left" valign="top">Mental health + suicidality</td><td align="left" valign="top">0.84</td><td align="left" valign="top">0.82-0.86</td><td align="left" valign="top">0.49 (0.36)</td><td align="left" valign="top">0.06</td><td align="left" valign="top">5.02 (599)</td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top">Mental health + suicidality + medical history</td><td align="left" valign="top">0.86</td><td align="left" valign="top">0.83-0.87</td><td align="left" valign="top">0.45 (0.36)</td><td align="left" valign="top">0.03</td><td align="left" valign="top">3.84 (599)</td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top">Mental health + suicidality + medical history + depression words</td><td align="left" valign="top">0.87</td><td align="left" valign="top">0.84-0.88</td><td align="left" valign="top">0.44 (0.35)</td><td align="left" valign="top">0.01</td><td align="left" valign="top">1.66 (599)</td><td align="left" valign="top">.049</td></tr></tbody></table><table-wrap-foot><fn id="table2fn1"><p><sup>a</sup><italic>r</italic>: correlation between the actual and predicted values.</p></fn><fn id="table2fn2"><p><sup>b</sup>Mean abs residual: mean absolute residual of actual and predicted values.</p></fn><fn id="table2fn3"><p><sup>c</sup>res: mean difference of absolute model residuals.</p></fn><fn id="table2fn4"><p><sup>d</sup>N/A: not applicable.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s3-3"><title>Model Validation</title><p>The final model achieved a correlation of <italic>r</italic>=0.82 with the best-estimate assessments, 95% CI 0.77-0.86 in the held-out set, exceeding the <italic>r</italic>=0.70 threshold for criterion validity (H1; <xref ref-type="table" rid="table3">Table 3</xref>). Notably, the 2 individual psychotherapists agreed with each other at <italic>r</italic>=0.78, meaning the model&#x2019;s alignment with the best-estimate (<italic>r</italic>=0.82) already exceeds the level of agreement between individual clinicians. Correcting for imperfect interrater reliability among the psychotherapists increased the correlation to nearly perfect (<italic>r</italic><sub>corrected</sub>=0.96 with Krippendorff <italic><inline-formula><mml:math id="ieqn3"><mml:mi>&#x03B1;</mml:mi></mml:math></inline-formula></italic><sub>psychotherapists</sub> =0.73). This upper-bound estimate indicates that the observed <italic>r</italic>=0.82 is close to the maximum achievable given the reliability of the reference standard. Further, the correlation between the LBAs and student assessments was the same in the validation set as in the development set (<italic>r</italic>=0.87), indicating excellent generalization. The LBAs correlated more strongly with student assessments (<italic>r</italic>=0.87) than the best-estimate assessments did (<italic>r</italic>=0.82), likely because both the model and the students based their assessments on the same language responses.</p><table-wrap id="t3" position="float"><label>Table 3.</label><caption><p>Criterion validity: correlations with best-estimate assessments in validation set (n=212). The comprehensive LBA<sup><xref ref-type="table-fn" rid="table3fn1">a</xref></sup> model is trained on the mental health text, suicidality text, medical history text, and depression word variables; the parsimonious LBA model is trained solely on the mental health text. The third psychotherapist only assessed 101 participants and was therefore excluded from this table.</p></caption><table id="table3" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">1</td><td align="left" valign="bottom">2</td><td align="left" valign="bottom">3</td><td align="left" valign="bottom">4</td><td align="left" valign="bottom">5</td><td align="left" valign="bottom">6</td></tr></thead><tbody><tr><td align="left" valign="top">Comprehensive LBA</td><td align="left" valign="top">&#x2014;<sup><xref ref-type="table-fn" rid="table3fn2">b</xref></sup></td><td align="left" valign="top">.86<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.82<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.73<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.87<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top">Parsimonious LBA</td><td align="left" valign="top">.86<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">&#x2014;</td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.72<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.68<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.76<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top">Best-estimate assessment</td><td align="left" valign="top">.82<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">&#x2014;</td><td align="left" valign="top">.94<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.93<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.82<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top">Psychotherapist 1</td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.72<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.94<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">&#x2014;</td><td align="left" valign="top">.78<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top">Psychotherapist 2</td><td align="left" valign="top">.73<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.68<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.93<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.78<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">&#x2014;</td><td align="left" valign="top">.72<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td></tr><tr><td align="left" valign="top">Average student assessment</td><td align="left" valign="top">.87<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.76<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.82<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.77<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">.72<sup><xref ref-type="table-fn" rid="table3fn3">c</xref></sup></td><td align="left" valign="top">&#x2014;</td></tr></tbody></table><table-wrap-foot><fn id="table3fn1"><p><sup>a</sup>LBA: language-based assessment.</p></fn><fn id="table3fn2"><p><sup>b</sup>Not available.</p></fn><fn id="table3fn3"><p><sup>c</sup><italic>P</italic>&#x003C;.001 (2-sided test).</p></fn></table-wrap-foot></table-wrap><p>The LBAs further demonstrated strong convergent validity (H2; <xref ref-type="table" rid="table4">Table 4</xref>), including strong positive correlations with depression (<italic>r</italic><sub>PHQ</sub>=0.77), anxiety (<italic>r</italic><sub>GAD</sub>=0.72), stress (<italic>r</italic><sub>PSS</sub>=0.76), and suicidality scores (<italic>r</italic><sub>IDAS</sub>=0.62), as well as strong negative correlations with life satisfaction (<italic>r</italic><sub>SWLS</sub>=&#x2212;0.66) and harmony in life (<italic>r</italic><sub>HILS</sub>=&#x2212;0.64). All correlations were statistically significant (<italic>P</italic>&#x003C;.001) and classified as large effects [<xref ref-type="bibr" rid="ref77">77</xref>]. What is important to note is that the criterion and convergent validity estimates are not completely independent signals due to shared method variance: the rating scales were part of the data evaluated by the psychotherapists to achieve the best-estimate assessments, an inherent property of LEAD designs (ie, all data [<xref ref-type="bibr" rid="ref46">46</xref>]; see the Discussion section). Regarding external validity (H3), the LBAs correlated significantly with the number of health care visits (<italic>r</italic>=0.16, <italic>P</italic>&#x003C;.05), meeting the threshold for a small effect size [<xref ref-type="bibr" rid="ref77">77</xref>], but not with the number of sick days (<italic>r</italic>=0.09, not significant). Distributions of the LBAs and human assessments, as well as analyses of model residuals, are reported in Figures S2-S5 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>.</p><p>The parsimonious model using only mental health texts demonstrated strong performance, only modestly lower than the full model. The LBAs showed a correlation of <italic>r</italic>=0.79 with student assessments during model development (Table S10 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), and a correlation of <italic>r</italic>=0.77 with the best-estimate assessments in the validation set (<xref ref-type="table" rid="table3">Table 3</xref>). It showed significant large correlations with the PHQ-9, GAD-7, PSS-10, SWLS-3, and HILS-3, but a comparatively small correlation with IDAS suicidality scores (<xref ref-type="table" rid="table4">Table 4</xref> and Table S11 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). Correlations with the behavioral measures were in the expected direction, but did not reach significance.</p><p>In an exploratory post hoc analysis, we compared the LBAs against demographic, dictionary- and frequency-based baseline models trained and validated in the same pipeline. The LBAs outperformed all baselines, with the largest advantage for the parsimonious model in comparison to the language baseline models (see Figure S6 and Table S12 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). In a further exploratory analysis, the baseline LBAs predicted depression, anxiety, stress, and suicidality scores 10 weeks later, beyond the respective baseline scale score (Table S13 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>)</p><table-wrap id="t4" position="float"><label>Table 4.</label><caption><p>Convergent and external validity: correlations in validation set (n=212). The comprehensive LBA<sup><xref ref-type="table-fn" rid="table4fn1">a</xref></sup> model is trained on the mental health text, suicidality text, medical history text, and depression word variables; the parsimonious LBA model is trained solely on the mental health text. The <italic>P</italic> values of the correlations refer to a 2-sided test.</p></caption><table id="table4" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom" rowspan="2">Variable</td><td align="left" valign="bottom" colspan="4">LBA</td></tr><tr><td align="left" valign="bottom" colspan="2">Comprehensive</td><td align="left" valign="bottom" colspan="2">Parsimonious</td></tr><tr><td align="left" valign="top"/><td align="left" valign="top">Pearson <italic>r</italic></td><td align="left" valign="top"><italic>P</italic> value</td><td align="left" valign="top">Pearson <italic>r</italic></td><td align="left" valign="top"><italic>P</italic> value</td></tr></thead><tbody><tr><td align="left" valign="top">Convergent validity</td><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>PHQ-9<sup><xref ref-type="table-fn" rid="table4fn2">b</xref></sup></td><td align="left" valign="top">0.77</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">0.74<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>GAD-7<sup><xref ref-type="table-fn" rid="table4fn4">d</xref></sup></td><td align="left" valign="top">0.72</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">0.72<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>PSS-10<sup><xref ref-type="table-fn" rid="table4fn5">e</xref></sup></td><td align="left" valign="top">0.76</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">0.75<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>IDAS suicidality<sup><xref ref-type="table-fn" rid="table4fn6">f</xref></sup></td><td align="left" valign="top">0.62</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">0.46<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top">Inverse convergent validity</td><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>SWLS-3<sup><xref ref-type="table-fn" rid="table4fn7">g</xref></sup></td><td align="left" valign="top">&#x2212;0.66</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x2212;0.63<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>HILS-3<sup><xref ref-type="table-fn" rid="table4fn8">h</xref></sup></td><td align="left" valign="top">&#x2212;0.64</td><td align="left" valign="top">&#x003C;.001</td><td align="left" valign="top">&#x2212;0.64<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">&#x003C;.001</td></tr><tr><td align="left" valign="top">External validity</td><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/><td align="left" valign="top"/></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Sick days<sup><xref ref-type="table-fn" rid="table4fn9">i</xref></sup></td><td align="left" valign="top">0.09</td><td align="left" valign="top">.21</td><td align="left" valign="top">0.06</td><td align="left" valign="top">.42</td></tr><tr><td align="left" valign="top"><named-content content-type="indent">&#x00A0;&#x00A0;&#x00A0;&#x00A0;</named-content>Health care visits<sup><xref ref-type="table-fn" rid="table4fn10">j</xref></sup></td><td align="left" valign="top">0.16<sup><xref ref-type="table-fn" rid="table4fn3">c</xref></sup></td><td align="left" valign="top">.02</td><td align="left" valign="top">0.11</td><td align="left" valign="top">.11</td></tr></tbody></table><table-wrap-foot><fn id="table4fn1"><p><sup>a</sup>LBA: language-based assessment.</p></fn><fn id="table4fn2"><p><sup>b</sup>PHQ-9: Patient Health Questionnaire 9-item scale.</p></fn><fn id="table4fn3"><p><sup>c</sup>Update footnote.</p></fn><fn id="table4fn4"><p><sup>d</sup>GAD-7: Generalized Anxiety Disorder 7-item scale.</p></fn><fn id="table4fn5"><p><sup>e</sup>PSS-10: Perceived Stress Scale 10-item version.</p></fn><fn id="table4fn6"><p><sup>f</sup>IDAS suicidality: suicidality dimension of the Inventory of Depression and Anxiety Symptoms.</p></fn><fn id="table4fn7"><p><sup>g</sup>SWLS-3: Satisfaction With Life Scale 3-item version.</p></fn><fn id="table4fn8"><p><sup>h</sup>HILS-3: Harmony in Life Scale 3-item version.</p></fn><fn id="table4fn9"><p><sup>i</sup>Sick days: number of sick days taken in the last 3 months due to mental health issues.</p></fn><fn id="table4fn10"><p><sup>j</sup>Health care visits: number of health care visits taken in the last 3 months due to mental health issues.</p></fn></table-wrap-foot></table-wrap><p>Finally, the extracted LDA topics derived from the complete dataset (n=812) were well-interpretable for all 4 language responses (see <xref ref-type="fig" rid="figure4">Figure 4</xref>). We extracted 12 topics per text response variable and 2 topics per word response variable (see the Methods section). Eleven topics for mental health, 11 topics for suicidality, and 7 topics for medical history correlated significantly with the comprehensive LBA model. The significant topics ranged from |&#x03B2;|=0.08 to |&#x03B2;|=0.36. For the depression words, both topics were significantly associated with the LBAs, showing large associations of |&#x03B2;|=0.76. Overall, the number of significant topics exceeded the predefined threshold and aligned with established psychological theories [<xref ref-type="bibr" rid="ref49">49</xref>-<xref ref-type="bibr" rid="ref51">51</xref>], supporting face validity (H4).</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Topics associated with LBAs of high vs low need for mental health support (n=812). The horizontal scale shows the magnitude of the topic standardized regression weights (&#x03B2;), indicating the strength of the relationship between latent Dirichlet allocation (LDA) topics and the LBAs. The font color indicates their direction and significance: significantly negative (green), nonsignificant (gray), and significantly positive (red). Font size and transparency indicate the probability of a word within its topic. Topics were extracted and tested on the complete dataset (n=812), including the held-out validation set (see Methods for details). In the validation set, the topics&#x2018; association with the LBAs closely mirrored their associations with the best-estimate assessments (Figure S7 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). LBA: language-based assessment.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e105460_fig04.png"/></fig><p>Topics significantly associated with LBAs of low need for support included &#x201C;work,&#x201D; &#x201C;home,&#x201D; &#x201C;happy,&#x201D; &#x201C;fine,&#x201D; the absence of suicidal thoughts, good health, and rarely getting sick. Prominent words in the associated depression words topic included &#x201C;hopeful,&#x201D; &#x201C;content,&#x201D; and &#x201C;optimistic.&#x201D; These topics align with the broaden-and-build theory [<xref ref-type="bibr" rid="ref49">49</xref>] that suggests that positive emotions expand and facilitate cognitive and behavioral patterns, which increases well-being and social connectedness. Surprisingly, the mental health topics referring to anxiety, problems in life, and the COVID-19 restrictions were also significantly associated with lower need for mental health support, despite not being inherently positive.</p><p>Topics significantly associated with LBAs of high need for support referred to anxiety, depression, the perception of others (&#x201C;people&#x201D;); comparisons (&#x201C;like&#x201D;); thinking about, planning, or attempting suicide; and thinking that others would be better off without them. Within the medical history variable, both acute mental health issues (eg, taking psychiatric medication) and long-term mental health issues (indicated by words such as &#x201C;since,&#x201D; &#x201C;age,&#x201D; &#x201C;time,&#x201D; and &#x201C;suffered&#x201D;) were significantly related to higher LBAs. Prominent words in the associated depression words topic included &#x201C;anxiety,&#x201D; &#x201C;lethargic,&#x201D; and &#x201C;pessimistic.&#x201D; These topics correspond to theories of depression [<xref ref-type="bibr" rid="ref50">50</xref>,<xref ref-type="bibr" rid="ref51">51</xref>] that suggest that negative cognitive patterns and attentional biases lead to negative affect, which can lead to or reinforce symptoms of depression and anxiety.</p><p>The exploratory supervised embedding projections (in contrast to the LDA topics, only based on the held-out set; <xref ref-type="fig" rid="figure5">Figure 5</xref>) revealed many words significantly associated with LBAs of low or high need for mental health support. Overall, the direction and significance of words matched those found in the LDA topics (<xref ref-type="fig" rid="figure4">Figure 4</xref>). For instance, words such as &#x201C;happy,&#x201D; &#x201C;good,&#x201D; and &#x201C;hopeful&#x201D; were significantly associated with low need for support, while &#x201C;anxiety,&#x201D; &#x201C;depression,&#x201D; and &#x201C;overwhelmed&#x201D; were associated with high need. In the suicidality variable, words explicitly denying &#x201C;suicidal&#x201D; &#x201C;thoughts&#x201D; (eg, &#x201C;not,&#x201D; &#x201C;never,&#x201D; and &#x201C;no&#x201D;) were strongly associated with low need, whereas words describing suicidal intent (eg, &#x201C;plans,&#x201D; &#x201C;attempting,&#x201D; and &#x201C;committing&#x201D;) were associated with high need. Notably, the word &#x201C;blood&#x201D; in the projection plot of the medical history variable likely refers to high blood pressure, as inferred from the LDA topics. This example shows how LDA topics and 3-grams can provide slightly more contextual grouping, while supervised embedding projections more directly reflect the LBA model&#x2019;s underlying process. Overall, these results further support the face validity of the LBAs by corroborating the LDA topic findings using a different method [<xref ref-type="bibr" rid="ref80">80</xref>].</p><fig position="float" id="figure5"><label>Figure 5.</label><caption><p>Words associated with LBAs of high vs low assessed need for mental health support (n=212). Each word is projected onto the semantic axis reflecting LBAs of low (left) vs high (right) need for mental health support. The boxes indicate the number of words significantly negatively correlated (green), not significantly correlated (gray), and significantly positively correlated (red) with the LBAs. The plots were created using supervised embedding projections [<xref ref-type="bibr" rid="ref72">72</xref>]. LBA: language-based assessment.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e105460_fig05.png"/></fig></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>This study developed and validated LBAs of need for mental health support using natural language responses. We validated 2 models: a parsimonious model using only a mental health text response, and a comprehensive model that additionally includes texts about suicidal thoughts and medical history, as well as depression-related words. Responding to all 4 language variables took participants less than 5 minutes in total.</p><p>Using the SEMP framework, models and hypotheses were preregistered before testing on a held-out validation set. The LBAs in the held-out set demonstrated strong criterion validity, correlating strongly with the best-estimate assessments of experienced psychotherapists who had access to longitudinal patient reports. The LBAs further showed strong convergent validity, correlating strongly with established measures of depression, anxiety, stress, and suicidality, and strongly negatively with established measures of well-being. Significant associations with linguistic topics were consistent with established psychological theories: low need for support assessments were associated with positive emotions, relationships, good health, and lack of suicidal thoughts, while high need for support assessments were linked to depression, anxiety, suicidal thoughts, and an ongoing history of mental illness. These associations were confirmed through 2 independent methods, LDA topic modeling and supervised embedding projections.</p></sec><sec id="s4-2"><title>Best-Estimate Reference Assessments</title><p>Much progress has been made in LLMs&#x2019; contextual language understanding and predictive accuracy [<xref ref-type="bibr" rid="ref16">16</xref>,<xref ref-type="bibr" rid="ref28">28</xref>]. However, for LLM-driven mental health assessments, the absence of a ground truth makes model training particularly complex. Best-estimate assessments are crucial for establishing a valid reference standard against which models can be evaluated, and averaging the assessments of multiple psychotherapists based on longitudinal clinical data represents one of the strongest available approaches [<xref ref-type="bibr" rid="ref47">47</xref>,<xref ref-type="bibr" rid="ref88">88</xref>]. It is important to note that the longitudinal clinical data evaluated by the psychotherapists also included the rating scales (ie, all data in the LEAD design [<xref ref-type="bibr" rid="ref46">46</xref>]), which makes the benchmarks for criterion and convergent validity not completely independent due to shared method variance.</p><p>The observed interrater reliability among the individual psychotherapists (Krippendorff &#x03B1;=0.73) highlights that even experienced experts sometimes disagree on mental health assessments. Given that this reliability sets a theoretical upper limit on expected validity [<xref ref-type="bibr" rid="ref30">30</xref>], it is notable that the LBAs achieved a correlation of <italic>r</italic>=0.82 with the clinicians&#x2019; best-estimate assessments. This suggests that the model captures a more stable signal than a single clinician, likely because it applies the same weighting of linguistic features consistently across cases, whereas individual clinicians may vary in what information they prioritize. However, the model lacks the clinical flexibility to recognize atypical presentations or contextual factors not captured in the text, which remains an important advantage of human judgment. This underscores the potential for a complementary approach, where the model provides a consistent baseline assessment that clinicians can then refine based on their expertise.</p><p>Remarkably, the assessments of the graduate students&#x2014;on which the model was trained&#x2014;converged closely with the best-estimate assessments, despite the students&#x2019; more limited clinical experience and their access to only the baseline language responses and demographics. Yet, an exploratory post hoc analysis indicated that the psychotherapists&#x2019; divergence from the student assessments was largely systematic rather than random. The difference between student and best-estimate assessments was substantially explained by the clinical rating scales that only the psychotherapists had seen (see Figure S8 and Table S14 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> for the full variance-reconstruction analysis). Still, the high overall convergence supports training LBA models on assessments from carefully instructed raters when expert capacity is limited, provided that the validation is conducted against an expert-based reference standard.</p></sec><sec id="s4-3"><title>Incremental Value of Multiple Probed Language Responses</title><p>While the parsimonious model already demonstrated a high accuracy, adding an explicit question about suicidality significantly improved model performance. The parsimonious model seemed to particularly miss information about suicidality, as evidenced by the correlation with IDAS suicidality substantially increasing from the parsimonious to the comprehensive model. This advantage could also be observed for the criterion validity, yet it is partly attributable to how the criterion itself was constructed: the best-estimate assessments were based on psychotherapists&#x2019; evaluations of longitudinal clinical data centering on internalizing disorders, self-harm, and suicidality. Moreover, the psychotherapists assessed participants&#x2019; need for mental health support immediately after assessing self-harm and suicidality [<xref ref-type="bibr" rid="ref45">45</xref>] (see Table S2 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), likely inflating the correlation between these assessments. Besides design characteristics, suicidal ideation and especially intention are evidently directly related to someone&#x2019;s need for mental health support. Suicidal ideation remains a highly stigmatized topic and is often not disclosed spontaneously [<xref ref-type="bibr" rid="ref89">89</xref>], highlighting the importance of explicitly prompting for suicidality to enhance the sensitivity of LBA models to this critical issue. Crucially, meta-analyses and reviews show that actively asking individuals about suicide does not appear to induce or increase suicidal ideation [<xref ref-type="bibr" rid="ref90">90</xref>,<xref ref-type="bibr" rid="ref91">91</xref>], affirming that explicit screening is safe.</p></sec><sec id="s4-4"><title>Comparison With Traditional Language Analysis Approaches</title><p>To situate the contribution of the embedding-based approach within more traditional language analysis methods, we compared the LBAs with 2 established alternatives trained in the same pipeline: a tf-idf (term frequency-inverse document frequency) bag-of-words model and a model trained on a negation-aware suicide dictionary and LIWC-22 (Linguistic Inquiry and Word Count) dimensions [<xref ref-type="bibr" rid="ref92">92</xref>]. The LBAs outperformed both baselines in the held-out set, while both language-based baseline models clearly outperformed the demographics-only model (see Figure S6 and Table S12 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref> for the full analysis). Notably, the advantage of the embeddings was modest when all 4 language variables were available, but substantial for the parsimonious models based only on the single mental health response. This shows that contextual embeddings maintain accuracy with brief responses, the scenario most relevant for large-scale screening.</p><p>Beyond these traditional baselines, our approach can be situated relative to generative encoder-decoder architectures increasingly used for LBAs by prompting an LLM to directly output a score from text. Zero-shot LLMs can approach human-rater accuracy but remain prompt-sensitive, performing best combined with document-tuned embedding-and-regression pipelines [<xref ref-type="bibr" rid="ref93">93</xref>]. We do not argue for one method over the other, but a document-tuned encoder paired with ridge regression is well suited to person-level psychological assessment: its regularization supports stable estimates under sample size constraints, the model weights yield deterministic outputs that remain stable over time, and document-tuned LLM embeddings have been repeatedly benchmarked as strong, efficient representations of psychological constructs from text [<xref ref-type="bibr" rid="ref94">94</xref>]. Practically, this architecture is also lightweight: with 335 million parameters, the model can be run on standard hardware and entirely on-site. In applied settings such as digital health care providers, sensitive free-text responses can be processed without transmitting them to external providers, at low computational cost.</p></sec><sec id="s4-5"><title>Limitations and Future Research</title><p>The LBAs correlated strongly with mental health scales, yet only partly with behavioral measures. The null finding for sick days may reflect the measure&#x2019;s limited variance in this sample (n=190 or 90% of validation-set participants reported 0 sick days over the past 3 months), rather than an insensitivity of the LBA to functional impairment.</p><p>Comparing the LBAs to rating scales requires caution, since the best-estimate assessments were based on both language responses and rating scale scores, making them difficult to use as an independent benchmark. The students based their assessments solely on language responses, so it is unsurprising that the LBAs correlated more strongly with student assessments (<italic>r</italic>=0.87) than the PHQ-9 did (<italic>r</italic>=0.75). Conversely, the psychotherapists had access to rating scales alongside language and clinical data, which likely explains why the best-estimate assessments correlated slightly more strongly with the PHQ-9 (<italic>r</italic>=0.84) than with the LBAs (<italic>r</italic>=0.82). These patterns reflect what information each assessor had access to, rather than the inherent superiority of either method. However, a recent investigation of the PHQ-9 showed that most respondents misinterpreted its instructions, basing responses on perceived bothersomeness rather than symptom frequency as intended [<xref ref-type="bibr" rid="ref95">95</xref>]&#x2014;a limitation potentially shared by similar scales. While this points to the scales being imperfect benchmarks for convergent validity, it reinforces the value of the criterion validity analysis against the best-estimate assessments as the primary validation approach. Beyond accuracy, LBAs offer unique advantages: they capture nuanced, real-life expressions of mental health and avoid the rigidity of rating scales [<xref ref-type="bibr" rid="ref16">16</xref>,<xref ref-type="bibr" rid="ref28">28</xref>]. Additionally, clinicians can review the written responses themselves alongside the numerical score, enabling a hybrid approach in which statistical evaluation and clinical judgment complement each other.</p><p>The model exhibited a regression-to-the-mean effect by systematically overestimating low need for support cases and underestimating high need for support cases (see Figures S2-S5 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>). This is common in regularized predictive models since they penalize coefficients to prevent overfitting, which shrinks predictions toward the mean of the criterion [<xref ref-type="bibr" rid="ref48">48</xref>]. This can be mitigated by distribution-based corrections, which realign the distribution of the predictions with that of the human assessments&#x2014;either by transforming the criterion before training [<xref ref-type="bibr" rid="ref48">48</xref>] or by remapping the predictions post hoc. Our demonstration (<xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) shows that post hoc transformations substantially reduce the overestimation of low-severity cases and underestimation of high-severity cases&#x2014;without lowering criterion validity (Figure S9 and Table S15 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>).</p><p>The need for mental health support assessed here is limited to internalizing symptoms and suicidality, and the model was trained and validated on a sample enriched for these conditions. While this coverage addresses the most common clinical needs worldwide&#x2014;depression and anxiety alone account for approximately 63% of mental health diagnoses [<xref ref-type="bibr" rid="ref13">13</xref>]&#x2014;the model&#x2019;s performance on related (eg, eating disorders and posttraumatic stress disorder) or less related conditions (eg, attention-deficit/hyperactivity disorders and psychotic disorders) is unknown. Additionally, while the student and psychotherapist assessors also had access to participants&#x2019; descriptions of their social situation and support system, it is uncertain whether the LBA model captures, for example, social determinants (eg, loneliness) of support need. More broadly, applying the model to populations with different demographics, cultural backgrounds, or base rates of mental health conditions may reduce its accuracy or introduce biases [<xref ref-type="bibr" rid="ref96">96</xref>]. This concern is consistent with prior work showing differential accuracy across racial or ethnic groups, such as regarding youth depression scores [<xref ref-type="bibr" rid="ref97">97</xref>] and language markers of depression [<xref ref-type="bibr" rid="ref98">98</xref>].</p><p>Lastly, this study assesses the need for mental health support on a 5-point scale, which has the advantage of being simple, interpretable, and comparable across individuals. Beyond the numerical output, the practical value of the scale lies in how the item formulations map onto a graded set of actions for those assessed. A complementary approach could involve using a generative LLM to produce personalized verbal support recommendations based on the individual&#x2019;s responses. For example, lower scores could point to appropriate self-directed resources (eg, psychoeducation and lifestyle guidance), while higher scores could surface potentially suitable clinician-led initiatives and available crisis helplines. Embedding this into a digital health system could let the LBA function as a triage layer. The nondiagnostic, support-framed output may also help address attitudinal barriers to help-seeking by lowering reluctance and making screening feel lower-stakes. Future work could compare engagement with &#x201C;need for support&#x201D; vs diagnostic framings, and whether this effect is stronger among populations with higher baseline stigma around mental health conditions.</p></sec><sec id="s4-6"><title>Implications</title><p>The model could be used in research settings to screen participants&#x2019; need for mental health support, for example, to stratify study groups or identify individuals who may benefit from intervention. Its short application time and strong alignment with best-estimate assessments make it practical for large-scale screening where comprehensive clinical evaluation is not feasible.</p><p>In real-world settings, many people affected by mental health issues do not recognize a need for treatment, which makes help in early identification, support, and intervention crucial [<xref ref-type="bibr" rid="ref2">2</xref>,<xref ref-type="bibr" rid="ref4">4</xref>,<xref ref-type="bibr" rid="ref7">7</xref>]. LBAs of the need for mental health support could complement existing screening tools and early intervention efforts. Possible applications include digital health platforms, informational websites on mood-related disorders, or health insurance portals, where individuals could receive preliminary guidance based on their responses. As the model assesses the need for support rather than psychiatric diagnoses, it may also help to reduce stigma surrounding mental health issues [<xref ref-type="bibr" rid="ref33">33</xref>], encouraging more people to seek support. Implementing such a screening system in a real-world setting would require rigorous calibration to ensure fair, unbiased, and ethical prioritization. This is especially important when the target group differs in demographics, cultural background, or mental health condition base rates from the current enriched sample.</p></sec><sec id="s4-7"><title>Conclusions</title><p>This study demonstrates that an LBA model based on natural language responses can closely align with experienced psychotherapists&#x2019; assessments of need for mental health support, using less than 5 minutes of respondent time. The model demonstrated strong criterion, convergent, and face validity against best-estimate reference assessments based on longitudinal data. By assessing the overall need for support based on individuals&#x2019; own descriptions of their mental health, rather than predicting single diagnostic labels, this approach may offer a less stigmatizing and potentially a more flexible alternative suited for early-stage mental health screening. Both the comprehensive and parsimonious models are openly available at the LBA model library [<xref ref-type="bibr" rid="ref54">54</xref>] to support further and independent validation.</p></sec></sec></body><back><ack><p>We want to thank the participants as well as the psychotherapist and student assessors. We also want to thank Prof Dr Anna-Lena Schubert and Wanja Hemmerich, MSc, for their input during earlier stages of this project. Generative AI (ChatGPT [OpenAI]; Claude [Anthropic PBC]) was used solely to generate and refine analysis code and refine the wording of this paper. It was not used to generate or alter substantive content or contribute to this study&#x2019;s interpretation.</p></ack><notes><sec><title>Funding</title><p>OK, KK, and VCE were supported by FORTE (STY-2022/0007; 2022&#x2010;01022); OK was also funded by Marianne och Marcus Wallenbergs stiftelse (MMW 2021.0058). Computations were enabled by resources provided by the Swedish National Infrastructure for Computing (SNIC) at Chalmers University of Technology, partially funded by the Swedish Research Council (2018&#x2010;05973).</p></sec></notes><fn-group><fn fn-type="con"><p>Conceptualization: OK, CW, KK, VCE, HAS</p><p>Data curation: KK, OK, CW</p><p>Formal analysis: CW</p><p>Funding acquisition: KK, OK</p><p>Investigation: CW, KK, OK, VCE</p><p>Methodology: CW, OK, KK, VV, HAS</p><p>Project administration: CW, VCE, KK, OK</p><p>Resources: KK, OK, HAS</p><p>Software: OK, HAS, CW</p><p>Supervision: OK, KK, HAS</p><p>Validation: CW, OK, KK</p><p>Visualization: CW</p><p>Writing - original draft: CW, VCE</p><p>Writing - review &#x0026; editing: CW, VCE, OK, VV, KK, HAS</p></fn><fn fn-type="conflict"><p>OK and KK have founded a start-up that assesses mental health issues by analyzing natural language responses with AI. The other authors declare no conflicts of interest.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">EU</term><def><p>European Union</p></def></def-item><def-item><term id="abb2">GAD</term><def><p>Generalized Anxiety Disorder</p></def></def-item><def-item><term id="abb3">GAD-7</term><def><p>Generalized Anxiety Disorder 7-Item Scale</p></def></def-item><def-item><term id="abb4">HILS-3</term><def><p>Harmony in Life Scale 3</p></def></def-item><def-item><term id="abb5">IDAS</term><def><p>Inventory of Depression and Anxiety Symptoms</p></def></def-item><def-item><term id="abb6">LBA</term><def><p>language-based assessment</p></def></def-item><def-item><term id="abb7">LDA</term><def><p>latent Dirichlet allocation</p></def></def-item><def-item><term id="abb8">LEAD</term><def><p>Longitudinal Expert All Data</p></def></def-item><def-item><term id="abb9">LIWC-22</term><def><p>Linguistic Inquiry and Word Count</p></def></def-item><def-item><term id="abb10">LLM</term><def><p>large language model</p></def></def-item><def-item><term id="abb11">MDD</term><def><p>major depressive disorder</p></def></def-item><def-item><term id="abb12">PHQ-9</term><def><p>Patient Health Questionnaire-9</p></def></def-item><def-item><term id="abb13">PSS-10</term><def><p>Perceived Stress Scale 10</p></def></def-item><def-item><term id="abb14">SEMP</term><def><p>Sequential Evaluation With Model Preregistration</p></def></def-item><def-item><term id="abb15">SWLS-3</term><def><p>Satisfaction With Life Scale 3</p></def></def-item><def-item><term id="abb16">tf-idf</term><def><p>term frequency-inverse document frequency</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="report"><person-group person-group-type="author"><collab>European Commission: Directorate-General for Health and Food Safety and Ipsos European Public Affairs, [Directorate-General for Health and Food Safety, Ipsos European Public Affairs]</collab></person-group><article-title>Mental health &#x2013; report</article-title><year>2023</year><access-date>2026-08-29</access-date><publisher-name>Publications Office of the European Union</publisher-name><comment><ext-link ext-link-type="uri" xlink:href="https://data.europa.eu/doi/10.2875/48999">https://data.europa.eu/doi/10.2875/48999</ext-link></comment></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Moitra</surname><given-names>M</given-names> </name><name name-style="western"><surname>Santomauro</surname><given-names>D</given-names> </name><name name-style="western"><surname>Collins</surname><given-names>PY</given-names> </name><etal/></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Hanlon</surname><given-names>C</given-names> </name></person-group><article-title>The global gap in treatment coverage for major depressive disorder in 84 countries from 2000&#x2013;2019: a systematic review and Bayesian meta-regression analysis</article-title><source>PLoS Med</source><year>2022</year><month>02</month><volume>19</volume><issue>2</issue><fpage>e1003901</fpage><pub-id pub-id-type="doi">10.1371/journal.pmed.1003901</pub-id><pub-id pub-id-type="medline">35167593</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Orozco</surname><given-names>R</given-names> </name><name name-style="western"><surname>Vigo</surname><given-names>D</given-names> </name><name name-style="western"><surname>Benjet</surname><given-names>C</given-names> </name><etal/></person-group><article-title>Barriers to treatment for mental disorders in six countries of the Americas: a regional report from the World Mental Health Surveys</article-title><source>J Affect Disord</source><year>2022</year><month>04</month><day>15</day><volume>303</volume><fpage>273</fpage><lpage>285</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2022.02.031</pub-id><pub-id pub-id-type="medline">35176342</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Olsson</surname><given-names>S</given-names> </name><name name-style="western"><surname>Hensing</surname><given-names>G</given-names> </name><name name-style="western"><surname>Burstr&#x00F6;m</surname><given-names>B</given-names> </name><name name-style="western"><surname>L&#x00F6;ve</surname><given-names>J</given-names> </name></person-group><article-title>Unmet need for mental healthcare in a population sample in Sweden: a cross-sectional study of inequalities based on gender, education, and country of birth</article-title><source>Community Ment Health J</source><year>2021</year><month>04</month><volume>57</volume><issue>3</issue><fpage>470</fpage><lpage>481</lpage><pub-id pub-id-type="doi">10.1007/s10597-020-00668-7</pub-id><pub-id pub-id-type="medline">32617737</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Katz</surname><given-names>SJ</given-names> </name><name name-style="western"><surname>Kessler</surname><given-names>RC</given-names> </name><name name-style="western"><surname>Frank</surname><given-names>RG</given-names> </name><name name-style="western"><surname>Leaf</surname><given-names>P</given-names> </name><name name-style="western"><surname>Lin</surname><given-names>E</given-names> </name><name name-style="western"><surname>Edlund</surname><given-names>M</given-names> </name></person-group><article-title>The use of outpatient mental health services in the United States and Ontario: the impact of mental morbidity and perceived need for care</article-title><source>Am J Public Health</source><year>1997</year><month>07</month><volume>87</volume><issue>7</issue><fpage>1136</fpage><lpage>1143</lpage><pub-id pub-id-type="doi">10.2105/ajph.87.7.1136</pub-id><pub-id pub-id-type="medline">9240103</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Meadows</surname><given-names>G</given-names> </name><name name-style="western"><surname>Harvey</surname><given-names>C</given-names> </name><name name-style="western"><surname>Fossey</surname><given-names>E</given-names> </name><name name-style="western"><surname>Burgess</surname><given-names>P</given-names> </name></person-group><article-title>Assessing perceived need for mental health care in a community survey: development of the Perceived Need for Care Questionnaire (PNCQ)</article-title><source>Soc Psychiatry Psychiatr Epidemiol</source><year>2000</year><month>09</month><volume>35</volume><issue>9</issue><fpage>427</fpage><lpage>435</lpage><pub-id pub-id-type="doi">10.1007/s001270050260</pub-id><pub-id pub-id-type="medline">11089671</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Sareen</surname><given-names>J</given-names> </name><name name-style="western"><surname>Stein</surname><given-names>MB</given-names> </name><name name-style="western"><surname>Campbell</surname><given-names>DW</given-names> </name><name name-style="western"><surname>Hassard</surname><given-names>T</given-names> </name><name name-style="western"><surname>Menec</surname><given-names>V</given-names> </name></person-group><article-title>The relation between perceived need for mental health treatment, DSM diagnosis, and quality of life: a Canadian population-based survey</article-title><source>Can J Psychiatry</source><year>2005</year><month>02</month><volume>50</volume><issue>2</issue><fpage>87</fpage><lpage>94</lpage><pub-id pub-id-type="doi">10.1177/070674370505000203</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Alegr&#x00ED;a</surname><given-names>M</given-names> </name><name name-style="western"><surname>NeMoyer</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bagu&#x00E9;</surname><given-names>IF</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Alvarez</surname><given-names>K</given-names> </name></person-group><article-title>Social determinants of mental health: where we are and where we need to go</article-title><source>Curr Psychiatry Rep</source><year>2018</year><month>09</month><day>17</day><volume>20</volume><issue>11</issue><fpage>95</fpage><pub-id pub-id-type="doi">10.1007/s11920-018-0969-9</pub-id><pub-id pub-id-type="medline">30221308</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Alon</surname><given-names>N</given-names> </name><name name-style="western"><surname>Macrynikola</surname><given-names>N</given-names> </name><name name-style="western"><surname>Jester</surname><given-names>DJ</given-names> </name><etal/></person-group><article-title>Social determinants of mental health in major depressive disorder: umbrella review of 26 meta-analyses and systematic reviews</article-title><source>Psychiatry Res</source><year>2024</year><month>05</month><volume>335</volume><fpage>115854</fpage><pub-id pub-id-type="doi">10.1016/j.psychres.2024.115854</pub-id><pub-id pub-id-type="medline">38554496</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bolton</surname><given-names>D</given-names> </name></person-group><article-title>A revitalized biopsychosocial model: core theory, research paradigms, and clinical implications</article-title><source>Psychol Med</source><year>2023</year><month>12</month><volume>53</volume><issue>16</issue><fpage>7504</fpage><lpage>7511</lpage><pub-id pub-id-type="doi">10.1017/S0033291723002660</pub-id><pub-id pub-id-type="medline">37681273</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Basterfield</surname><given-names>C</given-names> </name><name name-style="western"><surname>Fitzsimmons-Craft</surname><given-names>EE</given-names> </name><name name-style="western"><surname>Taylor</surname><given-names>CB</given-names> </name><name name-style="western"><surname>Eisenberg</surname><given-names>D</given-names> </name><name name-style="western"><surname>Wilfley</surname><given-names>DE</given-names> </name><name name-style="western"><surname>Newman</surname><given-names>MG</given-names> </name></person-group><article-title>Internalizing psychopathology and its links to suicidal ideation, dysfunctional attitudes, and help-seeking readiness in a national sample of college students</article-title><source>J Affect Disord</source><year>2024</year><month>04</month><day>1</day><volume>350</volume><fpage>255</fpage><lpage>263</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2024.01.058</pub-id><pub-id pub-id-type="medline">38224742</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kotov</surname><given-names>R</given-names> </name><name name-style="western"><surname>Krueger</surname><given-names>RF</given-names> </name><name name-style="western"><surname>Watson</surname><given-names>D</given-names> </name><etal/></person-group><article-title>The Hierarchical Taxonomy of Psychopathology (HiTOP): a dimensional alternative to traditional nosologies</article-title><source>J Abnorm Psychol</source><year>2017</year><volume>126</volume><issue>4</issue><fpage>454</fpage><lpage>477</lpage><pub-id pub-id-type="doi">10.1037/abn0000258</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhang</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>X</given-names> </name><name name-style="western"><surname>Wu</surname><given-names>S</given-names> </name><etal/></person-group><article-title>Global, regional and national burden of anxiety and depression disorders from 1990 to 2021, and forecasts up to 2040</article-title><source>J Affect Disord</source><year>2026</year><month>01</month><volume>393</volume><issue>Pt A</issue><fpage>120299</fpage><pub-id pub-id-type="doi">10.1016/j.jad.2025.120299</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Pennebaker</surname><given-names>JW</given-names> </name><name name-style="western"><surname>Mehl</surname><given-names>MR</given-names> </name><name name-style="western"><surname>Niederhoffer</surname><given-names>KG</given-names> </name></person-group><article-title>Psychological aspects of natural language use: our words, our selves</article-title><source>Annu Rev Psychol</source><year>2003</year><volume>54</volume><issue>1</issue><fpage>547</fpage><lpage>577</lpage><pub-id pub-id-type="doi">10.1146/annurev.psych.54.101601.145041</pub-id><pub-id pub-id-type="medline">12185209</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Boyd</surname><given-names>RL</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><article-title>Natural language analysis and the psychology of verbal behavior: the past, present, and future states of the field</article-title><source>J Lang Soc Psychol</source><year>2021</year><month>01</month><volume>40</volume><issue>1</issue><fpage>21</fpage><lpage>41</lpage><pub-id pub-id-type="doi">10.1177/0261927x20967028</pub-id><pub-id pub-id-type="medline">34413563</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mihalcea</surname><given-names>R</given-names> </name><name name-style="western"><surname>Biester</surname><given-names>L</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>RL</given-names> </name><etal/></person-group><article-title>How developments in natural language processing help us in understanding human behaviour</article-title><source>Nat Hum Behav</source><year>2024</year><month>10</month><volume>8</volume><issue>10</issue><fpage>1877</fpage><lpage>1889</lpage><pub-id pub-id-type="doi">10.1038/s41562-024-01938-0</pub-id><pub-id pub-id-type="medline">39438680</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Vaswani</surname><given-names>A</given-names> </name><name name-style="western"><surname>Shazeer</surname><given-names>N</given-names> </name><name name-style="western"><surname>Parmar</surname><given-names>N</given-names> </name><etal/></person-group><article-title>Attention is all you need</article-title><source>arXiv</source><comment>Preprint posted online on  Aug 2, 2023</comment><pub-id pub-id-type="doi">10.48550/arXiv.1706.03762</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eichstaedt</surname><given-names>JC</given-names> </name><name name-style="western"><surname>Smith</surname><given-names>RJ</given-names> </name><name name-style="western"><surname>Merchant</surname><given-names>RM</given-names> </name><etal/></person-group><article-title>Facebook language predicts depression in medical records</article-title><source>Proc Natl Acad Sci USA</source><year>2018</year><month>10</month><day>30</day><volume>115</volume><issue>44</issue><fpage>11203</fpage><lpage>11208</lpage><pub-id pub-id-type="doi">10.1073/pnas.1802331115</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Entwistle</surname><given-names>C</given-names> </name><name name-style="western"><surname>Hoemann</surname><given-names>K</given-names> </name><name name-style="western"><surname>Nightingale</surname><given-names>SJ</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>RL</given-names> </name></person-group><article-title>Psychosocial dynamics of suicidality and nonsuicidal self-injury: a digital linguistic perspective</article-title><source>npj Ment Health Res</source><year>2025</year><month>07</month><day>8</day><volume>4</volume><issue>1</issue><fpage>28</fpage><pub-id pub-id-type="doi">10.1038/s44184-025-00142-w</pub-id><pub-id pub-id-type="medline">40629100</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hur</surname><given-names>JK</given-names> </name><name name-style="western"><surname>Heffner</surname><given-names>J</given-names> </name><name name-style="western"><surname>Feng</surname><given-names>GW</given-names> </name><name name-style="western"><surname>Joormann</surname><given-names>J</given-names> </name><name name-style="western"><surname>Rutledge</surname><given-names>RB</given-names> </name></person-group><article-title>Language sentiment predicts changes in depressive symptoms</article-title><source>Proc Natl Acad Sci USA</source><year>2024</year><month>09</month><day>24</day><volume>121</volume><issue>39</issue><fpage>e2321321121</fpage><pub-id pub-id-type="doi">10.1073/pnas.2321321121</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Vine</surname><given-names>V</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>RL</given-names> </name><name name-style="western"><surname>Pennebaker</surname><given-names>JW</given-names> </name></person-group><article-title>Natural emotion vocabularies as windows on distress and well-being</article-title><source>Nat Commun</source><year>2020</year><month>09</month><day>10</day><volume>11</volume><issue>1</issue><fpage>4525</fpage><pub-id pub-id-type="doi">10.1038/s41467-020-18349-0</pub-id><pub-id pub-id-type="medline">32913209</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Stamatis</surname><given-names>CA</given-names> </name><name name-style="western"><surname>Meyerhoff</surname><given-names>J</given-names> </name><name name-style="western"><surname>Liu</surname><given-names>T</given-names> </name><etal/></person-group><article-title>Prospective associations of text-message-based sentiment with symptoms of depression, generalized anxiety, and social anxiety</article-title><source>Depression Anxiety</source><year>2022</year><month>12</month><volume>39</volume><issue>12</issue><fpage>794</fpage><lpage>804</lpage><pub-id pub-id-type="doi">10.1002/da.23286</pub-id><pub-id pub-id-type="medline">36281621</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hull</surname><given-names>TD</given-names> </name><name name-style="western"><surname>Malgaroli</surname><given-names>M</given-names> </name><name name-style="western"><surname>Connolly</surname><given-names>PS</given-names> </name><name name-style="western"><surname>Feuerstein</surname><given-names>S</given-names> </name><name name-style="western"><surname>Simon</surname><given-names>NM</given-names> </name></person-group><article-title>Two-way messaging therapy for depression and anxiety: longitudinal response trajectories</article-title><source>BMC Psychiatry</source><year>2020</year><month>06</month><day>12</day><volume>20</volume><issue>1</issue><fpage>297</fpage><pub-id pub-id-type="doi">10.1186/s12888-020-02721-x</pub-id><pub-id pub-id-type="medline">32532225</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Nook</surname><given-names>EC</given-names> </name><name name-style="western"><surname>Hull</surname><given-names>TD</given-names> </name><name name-style="western"><surname>Nock</surname><given-names>MK</given-names> </name><name name-style="western"><surname>Somerville</surname><given-names>LH</given-names> </name></person-group><article-title>Linguistic measures of psychological distance track symptom levels and treatment outcomes in a large set of psychotherapy transcripts</article-title><source>Proc Natl Acad Sci USA</source><year>2022</year><month>03</month><day>29</day><volume>119</volume><issue>13</issue><fpage>e2114737119</fpage><pub-id pub-id-type="doi">10.1073/pnas.2114737119</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Malgaroli</surname><given-names>M</given-names> </name><name name-style="western"><surname>Hull</surname><given-names>TD</given-names> </name><name name-style="western"><surname>Zech</surname><given-names>JM</given-names> </name><name name-style="western"><surname>Althoff</surname><given-names>T</given-names> </name></person-group><article-title>Natural language processing for mental health interventions: a systematic review and research framework</article-title><source>Transl Psychiatry</source><year>2023</year><month>10</month><day>6</day><volume>13</volume><issue>1</issue><fpage>309</fpage><pub-id pub-id-type="doi">10.1038/s41398-023-02592-2</pub-id><pub-id pub-id-type="medline">37798296</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Stade</surname><given-names>EC</given-names> </name><name name-style="western"><surname>Ungar</surname><given-names>L</given-names> </name><name name-style="western"><surname>Eichstaedt</surname><given-names>JC</given-names> </name><name name-style="western"><surname>Sherman</surname><given-names>G</given-names> </name><name name-style="western"><surname>Ruscio</surname><given-names>AM</given-names> </name></person-group><article-title>Depression and anxiety have distinct and overlapping language patterns: results from a clinical interview</article-title><source>J Psychopathol Clin Sci</source><year>2023</year><month>11</month><volume>132</volume><issue>8</issue><fpage>972</fpage><lpage>983</lpage><pub-id pub-id-type="doi">10.1037/abn0000850</pub-id><pub-id pub-id-type="medline">37471025</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Garcia</surname><given-names>D</given-names> </name><name name-style="western"><surname>Sikstr&#x00F6;m</surname><given-names>S</given-names> </name></person-group><article-title>Semantic measures: using natural language processing to measure, differentiate, and describe psychological constructs</article-title><source>Psychol Methods</source><year>2019</year><volume>24</volume><issue>1</issue><fpage>92</fpage><lpage>115</lpage><pub-id pub-id-type="doi">10.1037/met0000191</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><article-title>Beyond rating scales: with targeted evaluation, large language models are poised for psychological assessment</article-title><source>Psychiatry Res</source><year>2024</year><month>03</month><volume>333</volume><fpage>115667</fpage><pub-id pub-id-type="doi">10.1016/j.psychres.2023.115667</pub-id><pub-id pub-id-type="medline">38290286</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gu</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name></person-group><article-title>Natural language response formats for assessing depression and worry with large language models: a sequential evaluation with model pre-registration</article-title><source>Assessment</source><year>2026</year><month>09</month><volume>33</volume><issue>6</issue><fpage>927</fpage><lpage>953</lpage><pub-id pub-id-type="doi">10.1177/10731911251364022</pub-id><pub-id pub-id-type="medline">40974258</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Sikstr&#x00F6;m</surname><given-names>S</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><article-title>Natural language analyzed with AI-based transformers predict traditional subjective well-being measures approaching the theoretical upper limits in accuracy</article-title><source>Sci Rep</source><year>2022</year><month>03</month><day>10</day><volume>12</volume><issue>1</issue><fpage>3918</fpage><pub-id pub-id-type="doi">10.1038/s41598-022-07520-w</pub-id><pub-id pub-id-type="medline">35273198</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Nilsson</surname><given-names>A</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>R</given-names> </name><name name-style="western"><surname>Ganesan</surname><given-names>AV</given-names> </name><etal/></person-group><article-title>Language-based assessments for experienced well-being: accuracy and external validity across behaviors, traits, and states</article-title><source>PsyArXiv</source><comment>Preprint posted online on  Sep 18, 2025</comment><pub-id pub-id-type="doi">10.31234/osf.io/dgnaf</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ben-Zeev</surname><given-names>D</given-names> </name><name name-style="western"><surname>Young</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Corrigan</surname><given-names>PW</given-names> </name></person-group><article-title>DSM-V and the stigma of mental illness</article-title><source>J Ment Health</source><year>2010</year><month>08</month><volume>19</volume><issue>4</issue><fpage>318</fpage><lpage>327</lpage><pub-id pub-id-type="doi">10.3109/09638237.2010.492484</pub-id><pub-id pub-id-type="medline">20636112</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Clement</surname><given-names>S</given-names> </name><name name-style="western"><surname>Schauman</surname><given-names>O</given-names> </name><name name-style="western"><surname>Graham</surname><given-names>T</given-names> </name><etal/></person-group><article-title>What is the impact of mental health-related stigma on help-seeking? A systematic review of quantitative and qualitative studies</article-title><source>Psychol Med</source><year>2015</year><month>01</month><volume>45</volume><issue>1</issue><fpage>11</fpage><lpage>27</lpage><pub-id pub-id-type="doi">10.1017/S0033291714000129</pub-id><pub-id pub-id-type="medline">24569086</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bauer</surname><given-names>B</given-names> </name><name name-style="western"><surname>Norel</surname><given-names>R</given-names> </name><name name-style="western"><surname>Leow</surname><given-names>A</given-names> </name><name name-style="western"><surname>Rached</surname><given-names>ZA</given-names> </name><name name-style="western"><surname>Wen</surname><given-names>B</given-names> </name><name name-style="western"><surname>Cecchi</surname><given-names>G</given-names> </name></person-group><article-title>Using large language models to understand suicidality in a social media-based taxonomy of mental health disorders: linguistic analysis of Reddit posts</article-title><source>JMIR Ment Health</source><year>2024</year><month>05</month><day>16</day><volume>11</volume><fpage>e57234</fpage><pub-id pub-id-type="doi">10.2196/57234</pub-id><pub-id pub-id-type="medline">38771256</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><name name-style="western"><surname>Lahnala</surname><given-names>A</given-names> </name><name name-style="western"><surname>Vankudari</surname><given-names>S</given-names> </name><etal/></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Zirikly</surname><given-names>A</given-names> </name><name name-style="western"><surname>Yates</surname><given-names>A</given-names> </name><name name-style="western"><surname>Desmet</surname><given-names>B</given-names> </name><name name-style="western"><surname>Ireland</surname><given-names>M</given-names> </name><name name-style="western"><surname>Bedrick</surname><given-names>S</given-names> </name><name name-style="western"><surname>MacAvaney</surname><given-names>S</given-names> </name><name name-style="western"><surname>Bar</surname><given-names>K</given-names> </name><name name-style="western"><surname>Ophir</surname><given-names>Y</given-names> </name></person-group><article-title>Linking language-based distortion detection to mental health outcomes</article-title><conf-name>Proceedings of the 10th Workshop on Computational Linguistics and Clinical Psychology (CLPsych 2025)</conf-name><conf-date>May 3-4, 2025</conf-date><conf-loc>Albuquerque, New Mexico</conf-loc><fpage>62</fpage><lpage>68</lpage><pub-id pub-id-type="doi">10.18653/v1/2025.clpsych-1.5</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Jackson</surname><given-names>C</given-names> </name></person-group><article-title>The General Health Questionnaire</article-title><source>Occup Med</source><year>2006</year><volume>57</volume><issue>1</issue><fpage>79</fpage><lpage>79</lpage><pub-id pub-id-type="doi">10.1093/occmed/kql169</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Gianfrancesco</surname><given-names>MA</given-names> </name><name name-style="western"><surname>Tamang</surname><given-names>S</given-names> </name><name name-style="western"><surname>Yazdany</surname><given-names>J</given-names> </name><name name-style="western"><surname>Schmajuk</surname><given-names>G</given-names> </name></person-group><article-title>Potential biases in machine learning algorithms using electronic health record data</article-title><source>JAMA Intern Med</source><year>2018</year><month>11</month><day>1</day><volume>178</volume><issue>11</issue><fpage>1544</fpage><lpage>1547</lpage><pub-id pub-id-type="doi">10.1001/jamainternmed.2018.3763</pub-id><pub-id pub-id-type="medline">30128552</pub-id></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wright-Berryman</surname><given-names>J</given-names> </name><name name-style="western"><surname>Cohen</surname><given-names>J</given-names> </name><name name-style="western"><surname>Haq</surname><given-names>A</given-names> </name><name name-style="western"><surname>Black</surname><given-names>DP</given-names> </name><name name-style="western"><surname>Pease</surname><given-names>JL</given-names> </name></person-group><article-title>Virtually screening adults for depression, anxiety, and suicide risk using machine learning and language from an open-ended interview</article-title><source>Front Psychiatry</source><year>2023</year><volume>14</volume><fpage>1143175</fpage><pub-id pub-id-type="doi">10.3389/fpsyt.2023.1143175</pub-id><pub-id pub-id-type="medline">37377466</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Shin</surname><given-names>D</given-names> </name><name name-style="western"><surname>Kim</surname><given-names>K</given-names> </name><name name-style="western"><surname>Lee</surname><given-names>SB</given-names> </name><etal/></person-group><article-title>Detection of depression and suicide risk based on text from clinical interviews using machine learning: possibility of a new objective diagnostic marker</article-title><source>Front Psychiatry</source><year>2022</year><volume>13</volume><fpage>801301</fpage><pub-id pub-id-type="doi">10.3389/fpsyt.2022.801301</pub-id><pub-id pub-id-type="medline">35686182</pub-id></nlm-citation></ref><ref id="ref40"><label>40</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ben-Zeev</surname><given-names>D</given-names> </name><name name-style="western"><surname>Young</surname><given-names>MA</given-names> </name></person-group><article-title>Accuracy of hospitalized depressed patients&#x2019; and healthy controls&#x2019; retrospective symptom reports</article-title><source>J Nerv Ment Dis</source><year>2010</year><volume>198</volume><issue>4</issue><fpage>280</fpage><lpage>285</lpage><pub-id pub-id-type="doi">10.1097/NMD.0b013e3181d6141f</pub-id></nlm-citation></ref><ref id="ref41"><label>41</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bowes</surname><given-names>SM</given-names> </name><name name-style="western"><surname>Ammirati</surname><given-names>RJ</given-names> </name><name name-style="western"><surname>Costello</surname><given-names>TH</given-names> </name><name name-style="western"><surname>Basterfield</surname><given-names>C</given-names> </name><name name-style="western"><surname>Lilienfeld</surname><given-names>SO</given-names> </name></person-group><article-title>Cognitive biases, heuristics, and logical fallacies in clinical practice: a brief field guide for practicing clinicians and supervisors</article-title><source>Prof Psychol: Res Pract</source><year>2020</year><volume>51</volume><issue>5</issue><fpage>435</fpage><lpage>445</lpage><pub-id pub-id-type="doi">10.1037/pro0000309</pub-id></nlm-citation></ref><ref id="ref42"><label>42</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cronbach</surname><given-names>LJ</given-names> </name><name name-style="western"><surname>Meehl</surname><given-names>PE</given-names> </name></person-group><article-title>Construct validity in psychological tests</article-title><source>Psychol Bull</source><year>1955</year><month>07</month><volume>52</volume><issue>4</issue><fpage>281</fpage><lpage>302</lpage><pub-id pub-id-type="doi">10.1037/h0040957</pub-id><pub-id pub-id-type="medline">13245896</pub-id></nlm-citation></ref><ref id="ref43"><label>43</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Reitsma</surname><given-names>JB</given-names> </name><name name-style="western"><surname>Rutjes</surname><given-names>AWS</given-names> </name><name name-style="western"><surname>Khan</surname><given-names>KS</given-names> </name><name name-style="western"><surname>Coomarasamy</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bossuyt</surname><given-names>PM</given-names> </name></person-group><article-title>A review of solutions for diagnostic accuracy studies with an imperfect or missing reference standard</article-title><source>J Clin Epidemiol</source><year>2009</year><month>08</month><volume>62</volume><issue>8</issue><fpage>797</fpage><lpage>806</lpage><pub-id pub-id-type="doi">10.1016/j.jclinepi.2009.02.005</pub-id></nlm-citation></ref><ref id="ref44"><label>44</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chekroud</surname><given-names>AM</given-names> </name><name name-style="western"><surname>Hawrilenko</surname><given-names>M</given-names> </name><name name-style="western"><surname>Loho</surname><given-names>H</given-names> </name><etal/></person-group><article-title>Illusory generalizability of clinical prediction models</article-title><source>Science</source><year>2024</year><month>01</month><day>12</day><volume>383</volume><issue>6679</issue><fpage>164</fpage><lpage>167</lpage><pub-id pub-id-type="doi">10.1126/science.adg8538</pub-id></nlm-citation></ref><ref id="ref45"><label>45</label><nlm-citation citation-type="report"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Eijsbroek</surname><given-names>V</given-names> </name><name name-style="western"><surname>Wiebel</surname><given-names>C</given-names> </name><etal/></person-group><article-title>Validity of language-based assessments of depression and anxiety in comparison to best-estimate expert assessments: a sequential evaluation with model pre-registration</article-title></nlm-citation></ref><ref id="ref46"><label>46</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Spitzer</surname><given-names>RL</given-names> </name></person-group><article-title>Psychiatric diagnosis: are clinicians still necessary?</article-title><source>Compr Psychiatry</source><year>1983</year><volume>24</volume><issue>5</issue><fpage>399</fpage><lpage>411</lpage><pub-id pub-id-type="doi">10.1016/0010-440x(83)90032-9</pub-id><pub-id pub-id-type="medline">6354575</pub-id></nlm-citation></ref><ref id="ref47"><label>47</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Eijsbroek</surname><given-names>VC</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name><etal/></person-group><article-title>The LEADING guideline: reporting standards for expert panel, best-estimate diagnosis, and Longitudinal Expert All Data (LEAD) methods</article-title><source>Compr Psychiatry</source><year>2025</year><month>08</month><volume>141</volume><fpage>152603</fpage><pub-id pub-id-type="doi">10.1016/j.comppsych.2025.152603</pub-id><pub-id pub-id-type="medline">40479924</pub-id></nlm-citation></ref><ref id="ref48"><label>48</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Ganesan</surname><given-names>AV</given-names> </name><name name-style="western"><surname>Boyd</surname><given-names>R</given-names> </name></person-group><article-title>Demonstrating high validity of a new AI-language assessment of PTSD: a sequential evaluation with model pre-registration</article-title><source>PsyArXiv</source><comment>Preprint posted online on  Feb 11, 2026</comment><pub-id pub-id-type="doi">10.31234/osf.io/xw24e</pub-id></nlm-citation></ref><ref id="ref49"><label>49</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Fredrickson</surname><given-names>BL</given-names> </name></person-group><article-title>The role of positive emotions in positive psychology. The broaden-and-build theory of positive emotions</article-title><source>Am Psychol</source><year>2001</year><month>03</month><volume>56</volume><issue>3</issue><fpage>218</fpage><lpage>226</lpage><pub-id pub-id-type="doi">10.1037//0003-066x.56.3.218</pub-id><pub-id pub-id-type="medline">11315248</pub-id></nlm-citation></ref><ref id="ref50"><label>50</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Beck</surname><given-names>AT</given-names> </name></person-group><article-title>Thinking and depression. I. idiosyncratic content and cognitive distortions</article-title><source>Arch Gen Psychiatry</source><year>1963</year><month>10</month><volume>9</volume><issue>4</issue><fpage>324</fpage><lpage>333</lpage><pub-id pub-id-type="doi">10.1001/archpsyc.1963.01720160014002</pub-id><pub-id pub-id-type="medline">14045261</pub-id></nlm-citation></ref><ref id="ref51"><label>51</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Clark</surname><given-names>LA</given-names> </name><name name-style="western"><surname>Watson</surname><given-names>D</given-names> </name></person-group><article-title>Tripartite model of anxiety and depression: psychometric evidence and taxonomic implications</article-title><source>J Abnorm Psychol</source><year>1991</year><month>08</month><volume>100</volume><issue>3</issue><fpage>316</fpage><lpage>336</lpage><pub-id pub-id-type="doi">10.1037//0021-843x.100.3.316</pub-id><pub-id pub-id-type="medline">1918611</pub-id></nlm-citation></ref><ref id="ref52"><label>52</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Wiebel</surname><given-names>C</given-names> </name><name name-style="western"><surname>Eijsbroek</surname><given-names>V</given-names> </name><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><etal/></person-group><article-title>Mental health recommendations</article-title><source>OSF</source><year>2024</year><month>09</month><access-date>2026-06-24</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://osf.io/2jxr4">https://osf.io/2jxr4</ext-link></comment></nlm-citation></ref><ref id="ref53"><label>53</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Wiebel</surname><given-names>C</given-names> </name><name name-style="western"><surname>Hemmerich</surname><given-names>W</given-names> </name><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><etal/></person-group><article-title>Mental health recommendations - validation</article-title><source>OSF</source><year>2024</year><month>12</month><day>9</day><access-date>2026-06-24</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://osf.io/fc28k">https://osf.io/fc28k</ext-link></comment></nlm-citation></ref><ref id="ref54"><label>54</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Nilsson</surname><given-names>AH</given-names> </name><name name-style="western"><surname>Eijsbroek</surname><given-names>VC</given-names> </name><name name-style="western"><surname>Gu</surname><given-names>Z</given-names> </name><etal/></person-group><article-title>The language-based assessment model library: open model sharing for independent validation and broader applications</article-title><source>Adv Methods Pract Psychol Sci</source><year>2026</year><month>04</month><volume>9</volume><issue>2</issue><fpage>25152459261419036</fpage><pub-id pub-id-type="doi">10.1177/25152459261419036</pub-id></nlm-citation></ref><ref id="ref55"><label>55</label><nlm-citation citation-type="web"><source>Mental health recommendations - a Hugging Face Space by cwiebel</source><access-date>2026-06-24</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://huggingface.co/spaces/cwiebel/mental-health-recommendations">https://huggingface.co/spaces/cwiebel/mental-health-recommendations</ext-link></comment></nlm-citation></ref><ref id="ref56"><label>56</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Palan</surname><given-names>S</given-names> </name><name name-style="western"><surname>Schitter</surname><given-names>C</given-names> </name></person-group><article-title>Prolific.ac&#x2014;a subject pool for online experiments</article-title><source>J Behav Exper Finance</source><year>2018</year><month>03</month><volume>17</volume><fpage>22</fpage><lpage>27</lpage><pub-id pub-id-type="doi">10.1016/j.jbef.2017.12.004</pub-id></nlm-citation></ref><ref id="ref57"><label>57</label><nlm-citation citation-type="book"><source>DSM-5 Task Force, American Psychiatric Association Diagnostic and Statistical Manual of Mental Disorders: DSM-5</source><year>2013</year><publisher-name>American Psychiatric Association</publisher-name><pub-id pub-id-type="other">978-0-89042-555-8</pub-id></nlm-citation></ref><ref id="ref58"><label>58</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kroenke</surname><given-names>K</given-names> </name><name name-style="western"><surname>Spitzer</surname><given-names>RL</given-names> </name><name name-style="western"><surname>Williams</surname><given-names>JBW</given-names> </name></person-group><article-title>The PHQ-9: validity of a brief depression severity measure</article-title><source>J Gen Intern Med</source><year>2001</year><month>09</month><volume>16</volume><issue>9</issue><fpage>606</fpage><lpage>613</lpage><pub-id pub-id-type="doi">10.1046/j.1525-1497.2001.016009606.x</pub-id><pub-id pub-id-type="medline">11556941</pub-id></nlm-citation></ref><ref id="ref59"><label>59</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Spitzer</surname><given-names>RL</given-names> </name><name name-style="western"><surname>Kroenke</surname><given-names>K</given-names> </name><name name-style="western"><surname>Williams</surname><given-names>JBW</given-names> </name><name name-style="western"><surname>L&#x00F6;we</surname><given-names>B</given-names> </name></person-group><article-title>A brief measure for assessing generalized anxiety disorder: the GAD-7</article-title><source>Arch Intern Med</source><year>2006</year><month>05</month><day>22</day><volume>166</volume><issue>10</issue><fpage>1092</fpage><lpage>1097</lpage><pub-id pub-id-type="doi">10.1001/archinte.166.10.1092</pub-id><pub-id pub-id-type="medline">16717171</pub-id></nlm-citation></ref><ref id="ref60"><label>60</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cohen</surname><given-names>S</given-names> </name><name name-style="western"><surname>Kamarck</surname><given-names>T</given-names> </name><name name-style="western"><surname>Mermelstein</surname><given-names>R</given-names> </name></person-group><article-title>A global measure of perceived stress</article-title><source>J Health Soc Behav</source><year>1983</year><month>12</month><volume>24</volume><issue>4</issue><fpage>385</fpage><lpage>396</lpage><pub-id pub-id-type="doi">10.2307/2136404</pub-id><pub-id pub-id-type="medline">6668417</pub-id></nlm-citation></ref><ref id="ref61"><label>61</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Watson</surname><given-names>D</given-names> </name><name name-style="western"><surname>O&#x2019;Hara</surname><given-names>MW</given-names> </name><name name-style="western"><surname>Simms</surname><given-names>LJ</given-names> </name><etal/></person-group><article-title>Development and validation of the Inventory of Depression and Anxiety Symptoms (IDAS)</article-title><source>Psychol Assess</source><year>2007</year><volume>19</volume><issue>3</issue><fpage>253</fpage><lpage>268</lpage><pub-id pub-id-type="doi">10.1037/1040-3590.19.3.253</pub-id></nlm-citation></ref><ref id="ref62"><label>62</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Diener</surname><given-names>E</given-names> </name></person-group><article-title>Abbreviated three-item versions of the Satisfaction with Life Scale and the Harmony in Life Scale yield as strong psychometric properties as the original scales</article-title><source>J Pers Assess</source><year>2021</year><volume>103</volume><issue>2</issue><fpage>183</fpage><lpage>194</lpage><pub-id pub-id-type="doi">10.1080/00223891.2020.1737093</pub-id><pub-id pub-id-type="medline">32167788</pub-id></nlm-citation></ref><ref id="ref63"><label>63</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Diener</surname><given-names>E</given-names> </name><name name-style="western"><surname>Emmons</surname><given-names>RA</given-names> </name><name name-style="western"><surname>Larsen</surname><given-names>RJ</given-names> </name><name name-style="western"><surname>Griffin</surname><given-names>S</given-names> </name></person-group><article-title>The Satisfaction With Life Scale</article-title><source>J Pers Assess</source><year>1985</year><month>02</month><volume>49</volume><issue>1</issue><fpage>71</fpage><lpage>75</lpage><pub-id pub-id-type="doi">10.1207/s15327752jpa4901_13</pub-id><pub-id pub-id-type="medline">16367493</pub-id></nlm-citation></ref><ref id="ref64"><label>64</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><name name-style="western"><surname>Daukantait&#x0117;</surname><given-names>D</given-names> </name><name name-style="western"><surname>Hefferon</surname><given-names>K</given-names> </name><name name-style="western"><surname>Sikstr&#x00F6;m</surname><given-names>S</given-names> </name></person-group><article-title>The Harmony in Life Scale complements the Satisfaction with Life Scale: expanding the conceptualization of the cognitive component of subjective well-being</article-title><source>Soc Indic Res</source><year>2016</year><month>03</month><volume>126</volume><issue>2</issue><fpage>893</fpage><lpage>919</lpage><pub-id pub-id-type="doi">10.1007/s11205-015-0903-z</pub-id></nlm-citation></ref><ref id="ref65"><label>65</label><nlm-citation citation-type="web"><source>Samaritans</source><access-date>2026-09-09</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.samaritans.org/">https://www.samaritans.org/</ext-link></comment></nlm-citation></ref><ref id="ref66"><label>66</label><nlm-citation citation-type="web"><source>988 Lifeline</source><access-date>2026-09-09</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://988lifeline.org/">https://988lifeline.org/</ext-link></comment></nlm-citation></ref><ref id="ref67"><label>67</label><nlm-citation citation-type="web"><source>SoSci Survey</source><access-date>2026-09-09</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.soscisurvey.de/">https://www.soscisurvey.de/</ext-link></comment></nlm-citation></ref><ref id="ref68"><label>68</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Li</surname><given-names>X</given-names> </name><name name-style="western"><surname>Li</surname><given-names>J</given-names> </name></person-group><article-title>AnglE-optimized text embeddings</article-title><source>arXiv</source><comment>Preprint posted online on  Dec 31, 2024</comment><comment><ext-link ext-link-type="uri" xlink:href="https://arxiv.org/abs/2309.12871">https://arxiv.org/abs/2309.12871</ext-link></comment><pub-id pub-id-type="doi">10.48550/arXiv.2309.12871</pub-id></nlm-citation></ref><ref id="ref69"><label>69</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Muennighoff</surname><given-names>N</given-names> </name><name name-style="western"><surname>Tazi</surname><given-names>N</given-names> </name><name name-style="western"><surname>Magne</surname><given-names>L</given-names> </name><name name-style="western"><surname>Reimers</surname><given-names>N</given-names> </name></person-group><article-title>MTEB: Massive Text Embedding Benchmark</article-title><conf-name>Proceedings of the 17th Conference of the European Chapter of the Association for Computational Linguistics</conf-name><conf-date>May 2-6, 2023</conf-date><pub-id pub-id-type="doi">10.18653/v1/2023.eacl-main.148</pub-id></nlm-citation></ref><ref id="ref70"><label>70</label><nlm-citation citation-type="report"><person-group person-group-type="author"><name name-style="western"><surname>B&#x00E5;ng</surname><given-names>O</given-names> </name><name name-style="western"><surname>Gu</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Nilsson</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Language-based affect assessments capture experiment-induced changes beyond rating scales</article-title><comment><ext-link ext-link-type="uri" xlink:href="https://osf.io/preprints/psyarxiv/phgjn_v1">https://osf.io/preprints/psyarxiv/phgjn_v1</ext-link></comment><pub-id pub-id-type="doi">10.31234/osf.io/phgjn_v1</pub-id></nlm-citation></ref><ref id="ref71"><label>71</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Marker</surname><given-names>A</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><article-title>Evaluating document-tuned transformer representations for person-level mental health assessment</article-title><conf-name>Proceedings of the 10th Workshop on Computational Linguistics and Clinical Psychology (CLPsych 2026)</conf-name><conf-date>Jul 4, 2026</conf-date><pub-id pub-id-type="doi">10.18653/v1/2026.clpsych-1.14</pub-id></nlm-citation></ref><ref id="ref72"><label>72</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name><name name-style="western"><surname>Giorgi</surname><given-names>S</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><article-title>The text-package: an R-package for analyzing and visualizing human language using natural language processing and transformers</article-title><source>Psychol Methods</source><year>2023</year><month>12</month><volume>28</volume><issue>6</issue><fpage>1478</fpage><lpage>1498</lpage><pub-id pub-id-type="doi">10.1037/met0000542</pub-id><pub-id pub-id-type="medline">37126041</pub-id></nlm-citation></ref><ref id="ref73"><label>73</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Feuerriegel</surname><given-names>S</given-names> </name><name name-style="western"><surname>Barrie</surname><given-names>C</given-names> </name><name name-style="western"><surname>Crockett</surname><given-names>MJ</given-names> </name><etal/></person-group><article-title>A reporting checklist for large language models in behavioural science</article-title><source>Nat Hum Behav</source><year>2026</year><month>07</month><volume>10</volume><issue>7</issue><fpage>1182</fpage><lpage>1186</lpage><pub-id pub-id-type="doi">10.1038/s41562-026-02492-7</pub-id><pub-id pub-id-type="medline">42265331</pub-id></nlm-citation></ref><ref id="ref74"><label>74</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Kusupati</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bhatt</surname><given-names>G</given-names> </name><name name-style="western"><surname>Rege</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Matryoshka representation learning</article-title><source>arXiv</source><comment>Preprint posted online on  Feb 8, 2024</comment><pub-id pub-id-type="doi">10.48550/arXiv.2205.13147</pub-id><pub-id pub-id-type="medline">39032859</pub-id></nlm-citation></ref><ref id="ref75"><label>75</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Terwee</surname><given-names>CB</given-names> </name><name name-style="western"><surname>Bot</surname><given-names>SDM</given-names> </name><name name-style="western"><surname>de Boer</surname><given-names>MR</given-names> </name><etal/></person-group><article-title>Quality criteria were proposed for measurement properties of health status questionnaires</article-title><source>J Clin Epidemiol</source><year>2007</year><month>01</month><volume>60</volume><issue>1</issue><fpage>34</fpage><lpage>42</lpage><pub-id pub-id-type="doi">10.1016/j.jclinepi.2006.03.012</pub-id><pub-id pub-id-type="medline">17161752</pub-id></nlm-citation></ref><ref id="ref76"><label>76</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Spearman</surname><given-names>C</given-names> </name></person-group><article-title>The proof and measurement of association between two things</article-title><source>Am J Psychol</source><year>1904</year><month>01</month><volume>15</volume><issue>1</issue><fpage>72</fpage><lpage>101</lpage><pub-id pub-id-type="doi">10.2307/1412159</pub-id></nlm-citation></ref><ref id="ref77"><label>77</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Cohen</surname><given-names>J</given-names> </name></person-group><source>Statistical Power Analysis for the Behavioral Sciences</source><year>2013</year><edition>2</edition><publisher-name>Routledge</publisher-name><pub-id pub-id-type="doi">10.4324/9780203771587</pub-id></nlm-citation></ref><ref id="ref78"><label>78</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Blei</surname><given-names>DM</given-names> </name><name name-style="western"><surname>Ng</surname><given-names>AY</given-names> </name><name name-style="western"><surname>Jordan</surname><given-names>MI</given-names> </name></person-group><article-title>Latent Dirichlet Allocation</article-title><source>J Mach Learn Res</source><year>2003</year><access-date>2026-08-29</access-date><volume>3</volume><fpage>993</fpage><lpage>1022</lpage><comment><ext-link ext-link-type="uri" xlink:href="https://dl.acm.org/doi/10.5555/944919.944937">https://dl.acm.org/doi/10.5555/944919.944937</ext-link></comment></nlm-citation></ref><ref id="ref79"><label>79</label><nlm-citation citation-type="web"><article-title>Poweranalyse f&#x00FC;r korrelationen</article-title><source>StatistikGuru</source><access-date>2026-09-09</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://statistikguru.de/rechner/poweranalyse-korrelation.html">https://statistikguru.de/rechner/poweranalyse-korrelation.html</ext-link></comment></nlm-citation></ref><ref id="ref80"><label>80</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Eijsbroek</surname><given-names>VC</given-names> </name><name name-style="western"><surname>Nilsson</surname><given-names>A</given-names> </name><name name-style="western"><surname>Ackermann</surname><given-names>L</given-names> </name><etal/></person-group><article-title>Multiple methods for visualizing human language: a tutorial for social and behavioural scientists</article-title><source>PsyArXiv</source><access-date>2026-08-29</access-date><comment>Preprint posted online on  Apr 25, 2026</comment><comment><ext-link ext-link-type="uri" xlink:href="https://osf.io/preprints/psyarxiv/nxfvr_v3/">https://osf.io/preprints/psyarxiv/nxfvr_v3/</ext-link></comment></nlm-citation></ref><ref id="ref81"><label>81</label><nlm-citation citation-type="web"><source>Snowball</source><access-date>2026-08-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://snowballstem.org/">https://snowballstem.org/</ext-link></comment></nlm-citation></ref><ref id="ref82"><label>82</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Benjamini</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Hochberg</surname><given-names>Y</given-names> </name></person-group><article-title>Controlling the false discovery rate: a practical and powerful approach to multiple testing</article-title><source>J R Stat Soc Ser B</source><year>1995</year><month>01</month><day>1</day><volume>57</volume><issue>1</issue><fpage>289</fpage><lpage>300</lpage><pub-id pub-id-type="doi">10.1111/j.2517-6161.1995.tb02031.x</pub-id></nlm-citation></ref><ref id="ref83"><label>83</label><nlm-citation citation-type="web"><person-group person-group-type="author"><collab>R Core Team</collab></person-group><article-title>R foundation for statistical computing</article-title><source>R: a language environment for statistical computing</source><year>2023</year><access-date>2026-08-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.R-project.org/">https://www.R-project.org/</ext-link></comment></nlm-citation></ref><ref id="ref84"><label>84</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Ackermann</surname><given-names>L</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name></person-group><article-title>Theharmonylab/topics: topics (version v0909)</article-title><source>Zenodo</source><year>2024</year><month>05</month><day>9</day><access-date>2026-09-09</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.5281/zenodo.11165378">https://doi.org/10.5281/zenodo.11165378</ext-link></comment></nlm-citation></ref><ref id="ref85"><label>85</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Wickham</surname><given-names>H</given-names> </name><name name-style="western"><surname>Fran&#x00E7;ois</surname><given-names>R</given-names> </name><name name-style="western"><surname>Henry</surname><given-names>L</given-names> </name><name name-style="western"><surname>Muller</surname><given-names>K</given-names> </name></person-group><article-title>dplyr: a grammar of data manipulation</article-title><source>CRAN: Package dplyr</source><access-date>2026-08-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://CRAN.R-project.org/package=dplyr">https://CRAN.R-project.org/package=dplyr</ext-link></comment></nlm-citation></ref><ref id="ref86"><label>86</label><nlm-citation citation-type="web"><source>ggplot2: elegant graphics for data analysis</source><access-date>2026-08-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://ggplot2.tidyverse.org">https://ggplot2.tidyverse.org</ext-link></comment></nlm-citation></ref><ref id="ref87"><label>87</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Gamer</surname><given-names>M</given-names> </name><name name-style="western"><surname>Lemon</surname><given-names>J</given-names> </name><name name-style="western"><surname>Fellows</surname><given-names>I</given-names> </name><name name-style="western"><surname>Singh</surname><given-names>P</given-names> </name></person-group><article-title>Irr: various coefficients of interrater reliability and agreement</article-title><source>CRAN: Package irr</source><access-date>2026-08-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://CRAN.R-project.org/package=irr">https://CRAN.R-project.org/package=irr</ext-link></comment></nlm-citation></ref><ref id="ref88"><label>88</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bertens</surname><given-names>LCM</given-names> </name><name name-style="western"><surname>Broekhuizen</surname><given-names>BDL</given-names> </name><name name-style="western"><surname>Naaktgeboren</surname><given-names>CA</given-names> </name><etal/></person-group><article-title>Use of expert panels to define the reference standard in diagnostic research: a systematic review of published methods and reporting</article-title><source>PLoS Med</source><year>2013</year><month>10</month><day>15</day><volume>10</volume><issue>10</issue><fpage>e1001531</fpage><pub-id pub-id-type="doi">10.1371/journal.pmed.1001531</pub-id></nlm-citation></ref><ref id="ref89"><label>89</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hallford</surname><given-names>DJ</given-names> </name><name name-style="western"><surname>Rusanov</surname><given-names>D</given-names> </name><name name-style="western"><surname>Winestone</surname><given-names>B</given-names> </name><name name-style="western"><surname>Kaplan</surname><given-names>R</given-names> </name><name name-style="western"><surname>Fuller-Tyszkiewicz</surname><given-names>M</given-names> </name><name name-style="western"><surname>Melvin</surname><given-names>G</given-names> </name></person-group><article-title>Disclosure of suicidal ideation and behaviours: a systematic review and meta-analysis of prevalence</article-title><source>Clin Psychol Rev</source><year>2023</year><month>04</month><volume>101</volume><fpage>102272</fpage><pub-id pub-id-type="doi">10.1016/j.cpr.2023.102272</pub-id><pub-id pub-id-type="medline">37001469</pub-id></nlm-citation></ref><ref id="ref90"><label>90</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Dazzi</surname><given-names>T</given-names> </name><name name-style="western"><surname>Gribble</surname><given-names>R</given-names> </name><name name-style="western"><surname>Wessely</surname><given-names>S</given-names> </name><name name-style="western"><surname>Fear</surname><given-names>NT</given-names> </name></person-group><article-title>Does asking about suicide and related behaviours induce suicidal ideation? What is the evidence?</article-title><source>Psychol Med</source><year>2014</year><month>12</month><volume>44</volume><issue>16</issue><fpage>3361</fpage><lpage>3363</lpage><pub-id pub-id-type="doi">10.1017/S0033291714001299</pub-id></nlm-citation></ref><ref id="ref91"><label>91</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>DeCou</surname><given-names>CR</given-names> </name><name name-style="western"><surname>Schumann</surname><given-names>ME</given-names> </name></person-group><article-title>On the iatrogenic risk of assessing suicidality: a meta&#x2010;analysis</article-title><source>Suicide Life Threat Behav</source><year>2018</year><month>10</month><volume>48</volume><issue>5</issue><fpage>531</fpage><lpage>543</lpage><pub-id pub-id-type="doi">10.1111/sltb.12368</pub-id></nlm-citation></ref><ref id="ref92"><label>92</label><nlm-citation citation-type="report"><person-group person-group-type="author"><name name-style="western"><surname>Boyd</surname><given-names>RL</given-names> </name><name name-style="western"><surname>Ashokkumar</surname><given-names>A</given-names> </name><name name-style="western"><surname>Seraj</surname><given-names>S</given-names> </name><name name-style="western"><surname>Pennebaker</surname><given-names>JW</given-names> </name></person-group><article-title>The development and psychometric properties of LIWC-22</article-title><year>2022</year><publisher-name>The University of Texas at Austin</publisher-name><pub-id pub-id-type="doi">10.13140/RG.2.2.23890.43205</pub-id></nlm-citation></ref><ref id="ref93"><label>93</label><nlm-citation citation-type="other"><person-group person-group-type="author"><name name-style="western"><surname>Kaliosis</surname><given-names>P</given-names> </name><name name-style="western"><surname>Ganesan</surname><given-names>AV</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>ONE</given-names> </name><etal/></person-group><article-title>A systematic evaluation of large language models for PTSD severity estimation: the role of contextual knowledge and modeling strategies</article-title><source>arXiv</source><comment>Preprint posted online on  Jun 14, 2026</comment><pub-id pub-id-type="doi">10.21203/rs.3.rs-8376581/v1</pub-id></nlm-citation></ref><ref id="ref94"><label>94</label><nlm-citation citation-type="confproc"><person-group person-group-type="author"><name name-style="western"><surname>Marker</surname><given-names>A</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><name name-style="western"><surname>Schwartz</surname><given-names>HA</given-names> </name></person-group><person-group person-group-type="editor"><name name-style="western"><surname>Zirikly</surname><given-names>A</given-names> </name><name name-style="western"><surname>Bar</surname><given-names>K</given-names> </name><name name-style="western"><surname>MacAvaney</surname><given-names>S</given-names> </name><name name-style="western"><surname>Ireland</surname><given-names>M</given-names> </name><name name-style="western"><surname>Ophir</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Atzil-Slonim</surname><given-names>D</given-names> </name><name name-style="western"><surname>Varadarajan</surname><given-names>V</given-names> </name><name name-style="western"><surname>Bedrick</surname><given-names>S</given-names> </name><name name-style="western"><surname>Desmet</surname><given-names>B</given-names> </name></person-group><article-title>Evaluating document-tuned transformer representations for person-level mental health assessment</article-title><conf-name>Proceedings of the 10th Workshop on Computational Linguistics and Clinical Psychology (CLPsych 2026)</conf-name><conf-loc>San Diego, CA</conf-loc><fpage>178</fpage><lpage>187</lpage><pub-id pub-id-type="doi">10.18653/v1/2026.clpsych-1.14</pub-id></nlm-citation></ref><ref id="ref95"><label>95</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Panayiotou</surname><given-names>M</given-names> </name><name name-style="western"><surname>Razum</surname><given-names>J</given-names> </name><name name-style="western"><surname>Eisele</surname><given-names>G</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>SB</given-names> </name><name name-style="western"><surname>Fried</surname><given-names>EI</given-names> </name><name name-style="western"><surname>Cohen</surname><given-names>ZD</given-names> </name></person-group><article-title>Interpretation issues with the Patient Health Questionnaire instructions</article-title><source>JAMA Psychiatry</source><year>2026</year><month>04</month><day>1</day><volume>83</volume><issue>4</issue><fpage>399</fpage><lpage>403</lpage><pub-id pub-id-type="doi">10.1001/jamapsychiatry.2025.3796</pub-id><pub-id pub-id-type="medline">41405895</pub-id></nlm-citation></ref><ref id="ref96"><label>96</label><nlm-citation citation-type="report"><person-group person-group-type="author"><name name-style="western"><surname>Sparhuber</surname><given-names>M</given-names> </name><name name-style="western"><surname>Eijsbroek</surname><given-names>V</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>K</given-names> </name><name name-style="western"><surname>Giorgi</surname><given-names>S</given-names> </name><name name-style="western"><surname>Kjell</surname><given-names>O</given-names> </name></person-group><article-title>Evidence of limited biases in probed language-based assessments</article-title></nlm-citation></ref><ref id="ref97"><label>97</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Vaughn-Coaxum</surname><given-names>RA</given-names> </name><name name-style="western"><surname>Mair</surname><given-names>P</given-names> </name><name name-style="western"><surname>Weisz</surname><given-names>JR</given-names> </name></person-group><article-title>Racial/ethnic differences in youth depression indicators</article-title><source>Clin Psychol Sci</source><year>2016</year><month>03</month><volume>4</volume><issue>2</issue><fpage>239</fpage><lpage>253</lpage><pub-id pub-id-type="doi">10.1177/2167702615591768</pub-id></nlm-citation></ref><ref id="ref98"><label>98</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rai</surname><given-names>S</given-names> </name><name name-style="western"><surname>Stade</surname><given-names>EC</given-names> </name><name name-style="western"><surname>Giorgi</surname><given-names>S</given-names> </name><etal/></person-group><article-title>Key language markers of depression on social media depend on race</article-title><source>Proc Natl Acad Sci U S A</source><year>2024</year><month>04</month><day>2</day><volume>121</volume><issue>14</issue><fpage>e2319837121</fpage><pub-id pub-id-type="doi">10.1073/pnas.2319837121</pub-id><pub-id pub-id-type="medline">38530887</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Tables, figures, and exploratory analyses.</p><media xlink:href="mental_v13i1e105460_app1.docx" xlink:title="DOCX File, 8729 KB"/></supplementary-material></app-group></back></article>