<?xml version="1.0" encoding="UTF-8"?>
<resource xmlns="http://datacite.org/schema/kernel-4" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://datacite.org/schema/kernel-4 http://schema.datacite.org/meta/kernel-4.5/metadata.xsd">
  <identifier identifierType="DOI">10.60507/FK2/BK14F8</identifier>
  <creators>
    <creator>
      <creatorName nameType="Personal">Schilling, Julia</creatorName>
      <givenName>Julia</givenName>
      <familyName>Schilling</familyName>
      <nameIdentifier nameIdentifierScheme="ORCID" schemeURI="https://orcid.org">https://orcid.org/0009-0005-4105-1415</nameIdentifier>
      <affiliation>Universität Bonn</affiliation>
    </creator>
    <creator>
      <creatorName nameType="Personal">Fuchs, Robert</creatorName>
      <givenName>Robert</givenName>
      <familyName>Fuchs</familyName>
      <nameIdentifier nameIdentifierScheme="ORCID" schemeURI="https://orcid.org">https://orcid.org/0000-0001-7694-062X</nameIdentifier>
      <affiliation>Universität Bonn</affiliation>
    </creator>
  </creators>
  <titles>
    <title>Raw data and analysis code for the study “Project 2025 as a technocratic blueprint: A corpus-based linguistic analysis of conservative governance discourse”</title>
  </titles>
  <publisher>bonndata</publisher>
  <publicationYear>2025</publicationYear>
  <subjects>
    <subject>Arts and Humanities</subject>
  </subjects>
  <contributors>
    <contributor contributorType="ContactPerson">
      <contributorName nameType="Personal">Schilling, Julia</contributorName>
      <givenName>Julia</givenName>
      <familyName>Schilling</familyName>
      <affiliation>Universität Bonn</affiliation>
    </contributor>
  </contributors>
  <dates>
    <date dateType="Submitted">2025-11-05</date>
    <date dateType="Available">2025-11-11</date>
  </dates>
  <resourceType resourceTypeGeneral="Dataset"/>
  <relatedIdentifiers>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/AO0DXN</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/EAAWDF</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/7S5YMO</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/IZUWVA</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/E0US4A</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/WPIVQ7</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/4MRUKO</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/KFLETH</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/NYCNQ2</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/YR2EVG</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/IJ8QLP</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/SXTZG5</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/MTPXVR</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/YYXKMI</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/KFPU60</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/QGVBAT</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/PHSGX8</relatedIdentifier>
    <relatedIdentifier relationType="HasPart" relatedIdentifierType="DOI">10.60507/FK2/BK14F8/FRVM41</relatedIdentifier>
  </relatedIdentifiers>
  <sizes>
    <size>7110</size>
    <size>9106</size>
    <size>13103889</size>
    <size>20533</size>
    <size>7007</size>
    <size>1057511</size>
    <size>12827129</size>
    <size>1058835</size>
    <size>7038</size>
    <size>330359</size>
    <size>2310611</size>
    <size>916161</size>
    <size>26326</size>
    <size>846946</size>
    <size>2154423</size>
    <size>303972</size>
    <size>4200254</size>
    <size>9188</size>
  </sizes>
  <formats>
    <format>text/csv</format>
    <format>text/x-python-script</format>
    <format>text/csv</format>
    <format>text/x-python-script</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/x-r-notebook</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/csv</format>
    <format>text/markdown</format>
  </formats>
  <version>1.0</version>
  <rightsList>
    <rights rightsURI="info:eu-repo/semantics/openAccess"/>
    <rights rightsURI="http://creativecommons.org/licenses/by/4.0" xml:lang="en">Creative Commons Attribution 4.0 International License.</rights>
  </rightsList>
  <descriptions>
    <description descriptionType="Abstract">This repository contains all data and scripts used for the study “Project 2025 as a Technocratic Blueprint: A Corpus-Based Linguistic Analysis of Conservative Governance Discourse” (Schilling &amp;amp; Fuchs, 2025).
The study investigates the language of the Heritage Foundation’s Project 2025, a 900-page conservative policy blueprint, using methods from Corpus-Assisted Discourse Studies (CADS), Political Discourse Analysis (PDA), and psycholinguistic text analysis. The corpus includes Project 2025 and Democratic and Republican Party platforms (2016–2024).

The dataset includes:
Raw data (/data/raw_data/): full texts of Project 2025 and party platforms in CSV format.

Processed data (/data/raw_data/Project2025_lemmaPOS.csv): tokenized, lemmatized, and POS-tagged text.

Keyness results (/data/keyness/): unigram and bigram keyness calculations (log-likelihood, log-ratio).

Collocation results (/data/collocations/): top 10 adjective, noun, and verb collocates per node.

LIWC results (/data/liwc/): LIWC-22 category scores for each corpus.

Analysis scripts (/code/): R Markdown file (analysis_project2025.Rmd) and two Python scripts for collocation and keyness analysis.

All files are in UTF-8 plain-text format. The dataset contains no personal, sensitive, or proprietary data and derives entirely from publicly accessible political documents.</description>
    <description descriptionType="TechnicalInfo">spaCy, 3.7</description>
    <description descriptionType="TechnicalInfo">R, 4.4.1</description>
    <description descriptionType="TechnicalInfo">RStudio, 2024.09.0+375</description>
    <description descriptionType="TechnicalInfo">Python, 3.11.8</description>
    <description descriptionType="TechnicalInfo">LIWC, 22</description>
  </descriptions>
</resource>
