# robots.txt for emailverify.io # Last updated: May 2026 # Default rules — all crawlers welcome on public pages User-agent: * Allow: / Allow: /llms.txt Allow: /llms-full.txt # Block unnecessary WordPress paths on the blog Disallow: /blog/page/amp/ Disallow: /blog/feed Disallow: /blog/feed/ Disallow: /blog/*/feed Disallow: /blog/wp-admin/ Disallow: /blog/wp-includes/ Disallow: /blog/wp-login.php Disallow: /blog/wp-register.php Disallow: /blog/wp-content/plugins/ Disallow: /blog/wp-content/cache/ Disallow: /blog/wp-content/themes/ Disallow: /blog/wp-json/ Disallow: /blog/*?s= Disallow: /blog/search/ # Note: /blog/author/ is intentionally NOT blocked. # Author archive pages are part of our E-E-A-T signal and should be crawled. # Explicit allow rules for AI engine crawlers # These are redundant with the User-agent: * Allow: / rule above, # but stating them explicitly prevents WAF and security plugin # misconfigurations from accidentally blocking these bots. User-agent: GPTBot Allow: / User-agent: ChatGPT-User Allow: / User-agent: OAI-SearchBot Allow: / User-agent: ClaudeBot Allow: / User-agent: Claude-Web Allow: / User-agent: anthropic-ai Allow: / User-agent: PerplexityBot Allow: / User-agent: Perplexity-User Allow: / User-agent: Google-Extended Allow: / User-agent: Applebot-Extended Allow: / User-agent: Bytespider Allow: / User-agent: CCBot Allow: / User-agent: cohere-ai Allow: / User-agent: Meta-ExternalAgent Allow: / User-agent: Meta-ExternalFetcher Allow: / User-agent: Mistralai-User Allow: / User-agent: Diffbot Allow: / # Sitemaps Sitemap: https://www.emailverify.io/sitemap.xml