summaryrefslogtreecommitdiff
path: root/robots.txt
diff options
context:
space:
mode:
authorRadiohotline <radiohotline@disroot.org>2026-08-07 14:54:45 +0000
committerRadiohotline <radiohotline@disroot.org>2026-08-07 14:54:45 +0000
commite1baf35da18ae8a66fb32d4311c016afc999fdf9 (patch)
tree94ed32f8af92fc5b9a2b74891d037b31719a4fe4 /robots.txt
parent967effe5fe1aef78e37c068b0bb8979b2331dfc5 (diff)
article conclusionHEADmain
Diffstat (limited to 'robots.txt')
-rw-r--r--robots.txt48
1 files changed, 48 insertions, 0 deletions
diff --git a/robots.txt b/robots.txt
new file mode 100644
index 0000000..92f6a08
--- /dev/null
+++ b/robots.txt
@@ -0,0 +1,48 @@
+# This file tells search engines and bots what they are allowed to see on your site.
+
+# This is the default rule, which allows search engines to crawl your site (recommended).
+User-agent: *
+Allow: /
+
+# If you do not want AI bots to crawl your site, remove the # from the following lines:
+User-agent: AI2Bot
+User-agent: Ai2Bot-Dolma
+User-agent: Amazonbot
+User-agent: anthropic-ai
+User-agent: Applebot-Extended
+User-agent: Bytespider
+User-agent: CCBot
+User-agent: ChatGPT-User
+User-agent: Claude-Web
+User-agent: ClaudeBot
+User-agent: cohere-ai
+User-agent: Diffbot
+User-agent: DuckAssistBot
+User-agent: FacebookBot
+User-agent: FriendlyCrawler
+User-agent: Google-Extended
+User-agent: GoogleOther
+User-agent: GoogleOther-Image
+User-agent: GoogleOther-Video
+User-agent: GPTBot
+User-agent: iaskspider/2.0
+User-agent: ICC-Crawler
+User-agent: ImagesiftBot
+User-agent: img2dataset
+User-agent: ISSCyberRiskCrawler
+User-agent: Kangaroo Bot
+User-agent: Meta-ExternalAgent
+User-agent: Meta-ExternalFetcher
+User-agent: OAI-SearchBot
+User-agent: omgili
+User-agent: omgilibot
+User-agent: PanguBot
+User-agent: PerplexityBot
+User-agent: PetalBot
+User-agent: Scrapy
+User-agent: Sidetrade indexer bot
+User-agent: Timpibot
+User-agent: VelenPublicWebCrawler
+User-agent: Webzio-Extended
+User-agent: YouBot
+Disallow: /