{
  "entity_a": {
    "slug": "chatgpt-user",
    "name": "ChatGPT-User"
  },
  "entity_b": {
    "slug": "googleother",
    "name": "GoogleOther"
  },
  "fields_documented": 12,
  "claims": 18,
  "last_verified": "2026-09-15",
  "fields": [
    {
      "field": "cites_sources",
      "value_a": "not_documented",
      "value_b": null,
      "claim_id_a": 25,
      "claim_id_b": null,
      "same": false
    },
    {
      "field": "honours_crawl_delay",
      "value_a": "not_documented",
      "value_b": "no",
      "claim_id_a": 20,
      "claim_id_b": 72,
      "same": false
    },
    {
      "field": "observed_fetches_robots_txt",
      "value_a": null,
      "value_b": "no",
      "claim_id_a": null,
      "claim_id_b": 147,
      "same": false
    },
    {
      "field": "observed_requested_mcp_endpoint",
      "value_a": null,
      "value_b": "yes",
      "claim_id_a": null,
      "claim_id_b": 146,
      "same": false
    },
    {
      "field": "publishes_ip_list",
      "value_a": "yes",
      "value_b": "yes",
      "claim_id_a": 21,
      "claim_id_b": 73,
      "same": true
    },
    {
      "field": "purpose_documented",
      "value_a": "user_fetch",
      "value_b": "mixed",
      "claim_id_a": 23,
      "claim_id_b": 75,
      "same": false
    },
    {
      "field": "reads_llms_txt",
      "value_a": "not_documented",
      "value_b": null,
      "claim_id_a": 26,
      "claim_id_b": null,
      "same": false
    },
    {
      "field": "respects_robots_txt",
      "value_a": "partial",
      "value_b": "yes",
      "claim_id_a": 19,
      "claim_id_b": 71,
      "same": false
    },
    {
      "field": "robots_txt_token",
      "value_a": null,
      "value_b": "GoogleOther",
      "claim_id_a": null,
      "claim_id_b": 77,
      "same": false
    },
    {
      "field": "separate_search_and_training_tokens",
      "value_a": "not_documented",
      "value_b": "not_documented",
      "claim_id_a": 24,
      "claim_id_b": 76,
      "same": true
    },
    {
      "field": "user_agent_string_full",
      "value_a": "Mozilla/5.0 AppleWebKit/537.36 (KHTML, like Gecko); compatible; ChatGPT-User/1.0; +https://openai.com/bot",
      "value_b": null,
      "claim_id_a": 27,
      "claim_id_b": null,
      "same": false
    },
    {
      "field": "verifiable_by_rdns",
      "value_a": "not_documented",
      "value_b": "yes",
      "claim_id_a": 22,
      "claim_id_b": 74,
      "same": false
    }
  ],
  "differences": [
    {
      "field": "cites_sources",
      "statement_a": "OpenAI does not publish whether ChatGPT responses cite pages fetched by ChatGPT-User.",
      "statement_b": null,
      "claim_id_a": 25,
      "claim_id_b": null
    },
    {
      "field": "honours_crawl_delay",
      "statement_a": "OpenAI does not publish whether ChatGPT-User honours the robots.txt Crawl-delay directive.",
      "statement_b": "Google does not support the robots.txt Crawl-delay directive for GoogleOther.",
      "claim_id_a": 20,
      "claim_id_b": 72
    },
    {
      "field": "observed_fetches_robots_txt",
      "statement_a": null,
      "statement_b": "GoogleOther made 93 requests to Rattlesnakes By Mail between 2026-09-13T20:00Z and 2026-09-15T09:00Z and fetched /robots.txt on none of those requests. GoogleOther's first request to Rattlesnakes By Mail was the home page at 2026-09-14T12:03:30Z.",
      "claim_id_a": null,
      "claim_id_b": 147
    },
    {
      "field": "observed_requested_mcp_endpoint",
      "statement_a": null,
      "statement_b": "GoogleOther requested the /mcp endpoint of Rattlesnakes By Mail twice between 2026-09-13T20:00Z and 2026-09-15T09:00Z, at 2026-09-14T15:07:30Z and at 2026-09-14T18:27:49Z. Rattlesnakes By Mail returned status 405 to both GoogleOther requests for /mcp.",
      "claim_id_a": null,
      "claim_id_b": 146
    },
    {
      "field": "purpose_documented",
      "statement_a": "ChatGPT-User visits web pages for user actions in ChatGPT and in Custom GPTs.",
      "statement_b": "GoogleOther is a generic crawler that Google product teams use to fetch publicly accessible content from sites.",
      "claim_id_a": 23,
      "claim_id_b": 75
    },
    {
      "field": "reads_llms_txt",
      "statement_a": "OpenAI does not publish whether ChatGPT-User reads llms.txt files.",
      "statement_b": null,
      "claim_id_a": 26,
      "claim_id_b": null
    },
    {
      "field": "respects_robots_txt",
      "statement_a": "ChatGPT-User fetches pages in response to user actions in ChatGPT, and OpenAI states that robots.txt rules may not apply to ChatGPT-User fetches.",
      "statement_b": "GoogleOther obeys robots.txt rules when crawling automatically.",
      "claim_id_a": 19,
      "claim_id_b": 71
    },
    {
      "field": "robots_txt_token",
      "statement_a": null,
      "statement_b": "The robots.txt token for GoogleOther is GoogleOther.",
      "claim_id_a": null,
      "claim_id_b": 77
    },
    {
      "field": "user_agent_string_full",
      "statement_a": "ChatGPT-User sends the user agent string Mozilla/5.0 AppleWebKit/537.36 (KHTML, like Gecko); compatible; ChatGPT-User/1.0; +https://openai.com/bot.",
      "statement_b": null,
      "claim_id_a": 27,
      "claim_id_b": null
    },
    {
      "field": "verifiable_by_rdns",
      "statement_a": "OpenAI does not publish a reverse-DNS verification method for ChatGPT-User.",
      "statement_b": "GoogleOther requests are verifiable by reverse DNS lookup resolving to googlebot.com, google.com, or googleusercontent.com.",
      "claim_id_a": 22,
      "claim_id_b": 74
    }
  ],
  "sources": [
    "https://developers.google.com/crawling/docs/crawlers-fetchers/verifying-googlebot",
    "https://developers.google.com/search/blog/2019/07/a-note-on-unsupported-rules-in-robotstxt",
    "https://developers.google.com/search/docs/crawling-indexing/google-common-crawlers",
    "https://developers.google.com/static/crawling/ipranges/common-crawlers.json",
    "https://developers.openai.com/api/docs/bots",
    "https://openai.com/chatgpt-user.json",
    "https://rattlesnakesbymail.com/observed/googleother"
  ]
}