Search Options

Results per page
Sort
Preferred Languages
Advance

Results 1 - 10 of 22 for filemap (0.06 sec)

  1. src/main/java/org/codelibs/fess/dict/DictionaryManager.java

                            }
                        }
                    } catch (final Exception e) {
                        final String filePath = fileMap.get("path") != null ? fileMap.get("path").toString() : "unknown";
                        final String fileTimestamp = fileMap.get("@timestamp") != null ? fileMap.get("@timestamp").toString() : "unknown";
    Registered: Sat Dec 20 09:19:18 UTC 2025
    - Last Modified: Fri Nov 28 16:29:12 UTC 2025
    - 8K bytes
    - Viewed (0)
  2. fess-crawler/src/test/java/org/codelibs/fess/crawler/processor/impl/SitemapsResponseProcessorTest.java

            ResponseData responseData = new ResponseData();
            byte[] content = "<sitemap></sitemap>".getBytes();
            responseData.setResponseBody(content);
    
            SitemapUrl sitemap = new SitemapUrl();
            sitemap.setLoc("https://example.com/page1");
    
            SitemapSet sitemapSet = new SitemapSet();
            sitemapSet.addSitemap(sitemap);
    
            when(crawlerContainer.getComponent("sitemapsHelper")).thenReturn(sitemapsHelper);
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Thu Nov 13 13:29:22 UTC 2025
    - 12K bytes
    - Viewed (0)
  3. fess-crawler/src/main/java/org/codelibs/fess/crawler/entity/SitemapUrl.java

    /**
     * Represents a URL entry within a sitemap.
     *
     * <p>
     * This class encapsulates the properties of a URL as defined in the sitemap XML format,
     * including its location, last modification date, change frequency, and priority.
     * It also supports sitemap extensions such as images, videos, news, and alternate links.
     * It implements the {@link Sitemap} interface.
     * </p>
     *
     * <p>
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Thu Nov 13 13:34:36 UTC 2025
    - 9.1K bytes
    - Viewed (0)
  4. fess-crawler/src/test/java/org/codelibs/fess/crawler/helper/SitemapsHelperTest.java

                            + "    <loc>http://www.example.com/sitemap2.xml</loc>\n" + "  </sitemap>\n" + "</sitemapindex>";
            final InputStream in = new ByteArrayInputStream(xml.getBytes());
            final SitemapSet sitemapSet = sitemapsHelper.parse(in);
            final Sitemap[] sitemaps = sitemapSet.getSitemaps();
    
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Mon Nov 24 03:59:47 UTC 2025
    - 36.7K bytes
    - Viewed (0)
  5. fess-crawler/src/test/resources/org/codelibs/fess/crawler/helper/robots_malformed.txt

    User-agent: Bot1
    User-agent: Bot2
    User-agent: Bot3
    Disallow: /shared/
    
    # Case 11: Sitemap with various formats
    Sitemap: http://example.com/sitemap.xml
    sitemap: http://example.com/sitemap2.xml
    SITEMAP: http://example.com/sitemap3.xml
    Sitemap:    # empty sitemap (should be ignored)
    Sitemap: not-a-valid-url
    
    # Case 12: Malformed lines that should be completely ignored
    This line is completely invalid
    :NoKey
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Fri Nov 14 12:52:01 UTC 2025
    - 2.6K bytes
    - Viewed (0)
  6. fess-crawler/src/test/resources/org/codelibs/fess/crawler/helper/robots_wildcard.txt

    Allow: /page
    
    # Test multiple wildcards
    User-agent: MultiWildcardBot
    Disallow: /*.cgi*
    Disallow: /*?*id=*
    
    # Test literal $ in middle of pattern
    User-agent: DollarBot
    Disallow: /price$info
    
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Thu Nov 13 14:03:41 UTC 2025
    - 910 bytes
    - Viewed (0)
  7. fess-crawler/src/main/java/org/codelibs/fess/crawler/helper/SitemapsHelper.java

     */
    public class SitemapsHelper {
        private static final Logger logger = LogManager.getLogger(SitemapsHelper.class);
    
        /** The size of the preload buffer for checking file format. */
        protected int preloadSize = 512;
    
        /** Enable validation of sitemap entries. */
        protected boolean enableValidation = false;
    
        /** Maximum URL length according to sitemap specification. */
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Fri Nov 14 13:19:40 UTC 2025
    - 34.9K bytes
    - Viewed (0)
  8. fess-crawler/src/main/java/org/codelibs/fess/crawler/entity/RobotsTxt.java

        }
    
        /**
         * Adds a sitemap URL to the list of sitemaps.
         *
         * @param url The URL of the sitemap to be added
         */
        public void addSitemap(final String url) {
            if (!sitemapList.contains(url)) {
                sitemapList.add(url);
            }
        }
    
        /**
         * Returns an array of sitemap URLs.
         *
         * @return an array of sitemap URLs
         */
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Mon Nov 24 03:59:47 UTC 2025
    - 18.5K bytes
    - Viewed (0)
  9. fess-crawler/src/test/java/org/codelibs/fess/crawler/entity/RobotsTxtTest.java

            RobotsTxt robotsTxt = new RobotsTxt();
    
            robotsTxt.addSitemap("https://example.com/sitemap.xml");
            robotsTxt.addSitemap("https://example.com/sitemap2.xml");
    
            String[] sitemaps = robotsTxt.getSitemaps();
            assertEquals(2, sitemaps.length);
            assertEquals("https://example.com/sitemap.xml", sitemaps[0]);
            assertEquals("https://example.com/sitemap2.xml", sitemaps[1]);
        }
    
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Thu Nov 13 13:29:22 UTC 2025
    - 14.4K bytes
    - Viewed (0)
  10. fess-crawler/src/main/java/org/codelibs/fess/crawler/helper/RobotsTxtHelper.java

        protected static final Pattern CRAWL_DELAY_RECORD = Pattern.compile("^crawl-delay:\\s*([^\\s]+)\\s*$", Pattern.CASE_INSENSITIVE);
    
        /**
         * Pattern for Sitemap record.
         */
        protected static final Pattern SITEMAP_RECORD = Pattern.compile("^sitemap:\\s*([^\\s]+)\\s*$", Pattern.CASE_INSENSITIVE);
    
        /** Whether robots.txt processing is enabled. */
        protected boolean enabled = true;
    
        /**
    Registered: Sat Dec 20 11:21:39 UTC 2025
    - Last Modified: Fri Nov 14 12:52:01 UTC 2025
    - 11.4K bytes
    - Viewed (0)
Back to top