Lucene 4.8 - Facet Demo

package com.fox.facet;

/*

 * Licensed to the Apache Software Foundation (ASF) under one or more

 * contributor license agreements.  See the NOTICE file distributed with

 * this work for additional information regarding copyright ownership.

 * The ASF licenses this file to You under the Apache License, Version 2.0

 * (the "License"); you may not use this file except in compliance with

 * the License.  You may obtain a copy of the License at

 *

 *     http://www.apache.org/licenses/LICENSE-2.0

 *

 * Unless required by applicable law or agreed to in writing, software

 * distributed under the License is distributed on an "AS IS" BASIS,

 * WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.

 * See the License for the specific language governing permissions and

 * limitations under the License.

 */

import java.io.IOException;

import java.util.ArrayList;

import java.util.List;

import org.apache.lucene.analysis.core.WhitespaceAnalyzer;

import org.apache.lucene.document.Document;

import org.apache.lucene.facet.DrillDownQuery;

import org.apache.lucene.facet.DrillSideways;

import org.apache.lucene.facet.DrillSideways.DrillSidewaysResult;

import org.apache.lucene.facet.FacetField;

import org.apache.lucene.facet.FacetResult;

import org.apache.lucene.facet.Facets;

import org.apache.lucene.facet.FacetsCollector;

import org.apache.lucene.facet.FacetsConfig;

import org.apache.lucene.facet.taxonomy.FastTaxonomyFacetCounts;

import org.apache.lucene.facet.taxonomy.TaxonomyReader;

import org.apache.lucene.facet.taxonomy.directory.DirectoryTaxonomyReader;

import org.apache.lucene.facet.taxonomy.directory.DirectoryTaxonomyWriter;

import org.apache.lucene.index.DirectoryReader;

import org.apache.lucene.index.IndexWriter;

import org.apache.lucene.index.IndexWriterConfig;

import org.apache.lucene.search.IndexSearcher;

import org.apache.lucene.search.MatchAllDocsQuery;

import org.apache.lucene.store.Directory;

import org.apache.lucene.store.RAMDirectory;

import org.apache.lucene.util.Version;

/** Shows simple usage of faceted indexing and search. */

public class SimpleFacetsExample48 {

    private final Directory indexDir = new RAMDirectory();

    private final Directory taxoDir = new RAMDirectory();

    private final FacetsConfig config = new FacetsConfig();

    /** Empty constructor */

    public SimpleFacetsExample48() {

        config.setHierarchical("Publish Date", true);

    }

    /** Build the example index. */

    private void index() throws IOException {

        IndexWriter indexWriter = new IndexWriter(indexDir, new IndexWriterConfig(Version.LUCENE_48, new WhitespaceAnalyzer(

                Version.LUCENE_48)));

        // Writes facet ords to a separate directory from the main index

        DirectoryTaxonomyWriter taxoWriter = new DirectoryTaxonomyWriter(taxoDir);

        Document doc = new Document();

        doc.add(new FacetField("Author", "Bob"));

        doc.add(new FacetField("Publish Date", "2010", "10", "15"));

        indexWriter.addDocument(config.build(taxoWriter, doc));

        doc = new Document();

        doc.add(new FacetField("Author", "Lisa"));

        doc.add(new FacetField("Publish Date", "2010", "10", "20"));

        indexWriter.addDocument(config.build(taxoWriter, doc));

        doc = new Document();

        doc.add(new FacetField("Author", "Lisa"));

        doc.add(new FacetField("Publish Date", "2012", "1", "1"));

        indexWriter.addDocument(config.build(taxoWriter, doc));

        doc = new Document();

        doc.add(new FacetField("Author", "Susan"));

        doc.add(new FacetField("Publish Date", "2012", "1", "7"));

        indexWriter.addDocument(config.build(taxoWriter, doc));

        doc = new Document();

        doc.add(new FacetField("Author", "Frank"));

        doc.add(new FacetField("Publish Date", "1999", "5", "5"));

        indexWriter.addDocument(config.build(taxoWriter, doc));

        indexWriter.close();

        taxoWriter.close();

    }

    /** User runs a query and counts facets. */

    private List<FacetResult> facetsWithSearch() throws IOException {

        DirectoryReader indexReader = DirectoryReader.open(indexDir);

        IndexSearcher searcher = new IndexSearcher(indexReader);

        TaxonomyReader taxoReader = new DirectoryTaxonomyReader(taxoDir);

        FacetsCollector fc = new FacetsCollector();

        // MatchAllDocsQuery is for "browsing" (counts facets

        // for all non-deleted docs in the index); normally

        // you'd use a "normal" query:

        FacetsCollector.search(searcher, new MatchAllDocsQuery(), 10, fc);

        // Retrieve results

        List<FacetResult> results = new ArrayList<FacetResult>();

        // Count both "Publish Date" and "Author" dimensions

        Facets facets = new FastTaxonomyFacetCounts(taxoReader, config, fc);

        results.add(facets.getTopChildren(10, "Author"));

        results.add(facets.getTopChildren(10, "Publish Date"));

        indexReader.close();

        taxoReader.close();

        return results;

    }

    /**

     * User runs a query and counts facets only without collecting the matching

     * documents.

     */

    private List<FacetResult> facetsOnly() throws IOException {

        DirectoryReader indexReader = DirectoryReader.open(indexDir);

        IndexSearcher searcher = new IndexSearcher(indexReader);

        TaxonomyReader taxoReader = new DirectoryTaxonomyReader(taxoDir);

        FacetsCollector fc = new FacetsCollector();

        // MatchAllDocsQuery is for "browsing" (counts facets

        // for all non-deleted docs in the index); normally

        // you'd use a "normal" query:

        searcher.search(new MatchAllDocsQuery(), null /* Filter */, fc);

        // Retrieve results

        List<FacetResult> results = new ArrayList<FacetResult>();

        // Count both "Publish Date" and "Author" dimensions

        Facets facets = new FastTaxonomyFacetCounts(taxoReader, config, fc);

        results.add(facets.getTopChildren(10, "Author"));

        results.add(facets.getTopChildren(10, "Publish Date"));

        indexReader.close();

        taxoReader.close();

        return results;

    }

    /**

     * User drills down on 'Publish Date/2010', and we return facets for

     * 'Author'

     */

    private FacetResult drillDown() throws IOException {

        DirectoryReader indexReader = DirectoryReader.open(indexDir);

        IndexSearcher searcher = new IndexSearcher(indexReader);

        TaxonomyReader taxoReader = new DirectoryTaxonomyReader(taxoDir);

        // Passing no baseQuery means we drill down on all

        // documents ("browse only"):

        DrillDownQuery q = new DrillDownQuery(config);

        // Now user drills down on Publish Date/2010:

        q.add("Publish Date", "2010");

        FacetsCollector fc = new FacetsCollector();

        FacetsCollector.search(searcher, q, 10, fc);

        // Retrieve results

        Facets facets = new FastTaxonomyFacetCounts(taxoReader, config, fc);

        FacetResult result = facets.getTopChildren(10, "Author");

        indexReader.close();

        taxoReader.close();

        return result;

    }

    /**

     * User drills down on 'Publish Date/2010', and we return facets for both

     * 'Publish Date' and 'Author', using DrillSideways.

     */

    private List<FacetResult> drillSideways() throws IOException {

        DirectoryReader indexReader = DirectoryReader.open(indexDir);

        IndexSearcher searcher = new IndexSearcher(indexReader);

        TaxonomyReader taxoReader = new DirectoryTaxonomyReader(taxoDir);

        // Passing no baseQuery means we drill down on all

        // documents ("browse only"):

        DrillDownQuery q = new DrillDownQuery(config);

        // Now user drills down on Publish Date/2010:

        q.add("Publish Date", "2010");

        DrillSideways ds = new DrillSideways(searcher, config, taxoReader);

        DrillSidewaysResult result = ds.search(q, 10);

        // Retrieve results

        List<FacetResult> facets = result.facets.getAllDims(10);

        indexReader.close();

        taxoReader.close();

        return facets;

    }

    /** Runs the search example. */

    public List<FacetResult> runFacetOnly() throws IOException {

        index();

        return facetsOnly();

    }

    /** Runs the search example. */

    public List<FacetResult> runSearch() throws IOException {

        index();

        return facetsWithSearch();

    }

    /** Runs the drill-down example. */

    public FacetResult runDrillDown() throws IOException {

        index();

        return drillDown();

    }

    /** Runs the drill-sideways example. */

    public List<FacetResult> runDrillSideways() throws IOException {

        index();

        return drillSideways();

    }

    /** Runs the search and drill-down examples and prints the results. */

    public static void main(String[] args) throws Exception {

        System.out.println("Facet counting example:");

        System.out.println("-----------------------");

        SimpleFacetsExample48 example1 = new SimpleFacetsExample48();

        List<FacetResult> results1 = example1.runFacetOnly();

        System.out.println("Author: " + results1.get(0));

        System.out.println("Publish Date: " + results1.get(1));

        System.out.println("Facet counting example (combined facets and search):");

        System.out.println("-----------------------");

        SimpleFacetsExample48 example = new SimpleFacetsExample48();

        List<FacetResult> results = example.runSearch();

        System.out.println("Author: " + results.get(0));

        System.out.println("Publish Date: " + results.get(1));

        System.out.println("\n");

        System.out.println("Facet drill-down example (Publish Date/2010):");

        System.out.println("---------------------------------------------");

        System.out.println("Author: " + example.runDrillDown());

        System.out.println("\n");

        System.out.println("Facet drill-sideways example (Publish Date/2010):");

        System.out.println("---------------------------------------------");

        for (FacetResult result : example.runDrillSideways()) {

            System.out.println(result);

        }

    }

}

Result:

Facet counting example:

-----------------------

Author: dim=Author path=[] value=5 childCount=4

  Lisa (2)

  Bob (1)

  Susan (1)

  Frank (1)

Publish Date: dim=Publish Date path=[] value=5 childCount=3

  2010 (2)

  2012 (2)

  1999 (1)

Facet counting example (combined facets and search):

-----------------------

Author: dim=Author path=[] value=5 childCount=4

  Lisa (2)

  Bob (1)

  Susan (1)

  Frank (1)

Publish Date: dim=Publish Date path=[] value=5 childCount=3

  2010 (2)

  2012 (2)

  1999 (1)

Facet drill-down example (Publish Date/2010):

---------------------------------------------

Author: dim=Author path=[] value=4 childCount=2

  Bob (2)

  Lisa (2)

Facet drill-sideways example (Publish Date/2010):

---------------------------------------------

dim=Publish Date path=[] value=15 childCount=3

  2010 (6)

  2012 (6)

  1999 (3)

dim=Author path=[] value=6 childCount=2

  Bob (3)

  Lisa (3)

Lucene 4.8 - Facet Demo的更多相关文章

Lucene 4.3 - Facet demo
package com.fox.facet; import java.io.IOException; import java.util.ArrayList; import java.util.List ...
lucene 4.0 - Facet demo
package com.fox.facet; import java.io.File; import java.io.IOException; import java.util.ArrayList; ...
lucene搜索之facet查询原理和facet查询实例——TODO
转自:http://www.lai18.com/content/7084969.html Facet说明我们在浏览网站的时候,经常会遇到按某一类条件查询的情况,这种情况尤以电商网站最多,以天猫商城为 ...
（一）Lucene简介以及索引demo
一.百度百科 Lucene是apache软件基金会4 jakarta项目组的一个子项目,是一个开放源代码的全文检索引擎工具包,但它不是一个完整的全文检索引擎,而是一个全文检索引擎的架构,提供了完整的查 ...
Facet with Lucene
Facets with Lucene Posted on August 1, 2014 by Pascal Dimassimo in Latest Articles During the develo ...
MVC+MQ+WinServices+Lucene.Net Demo
前言: 我之前没有接触过Lucene.Net相关的知识,最近在园子里看到很多大神在分享这块的内容,深受启发.秉着“实践出真知”的精神,再结合公司项目的实际情况,有了写一个Demo的想法,算是对自己能力 ...
Lucene系列-facet
1.facet的直观认识 facet:面.切面.方面.个人理解就是维度,在满足query的前提下,观察结果在各维度上的分布(一个维度下各子类的数目). 如jd上搜“手机”,得到4009个商品.其中品牌 ...
lucene 4.4 demo
ackage com.zxf.demo; import java.io.BufferedReader; import java.io.File; import java.io.FileInputStr ...

随机推荐

HDU2027：统计元音
Problem Description 统计每个元音字母在字符串中出现的次数. Input 输入数据首先包括一个整数n,表示测试实例的个数,然后是n行长度不超过100的字符串. Output 对于每个 ...
opencv感兴趣区域ROI
addWeighted //显示原图 Mat src = imread("data/img/1.jpg"); imshow("src",src); //显示lo ...
炫龙笔记本的gtx965m显卡玩游戏很卡
这是我遇到的问题,我2016年10月份这样买了一款笔记本,主要看的是性价比吧!神舟.炫龙都是性价比,所以买了炫龙笔记本配置如下 cpu:i7 4870hq 显卡:gtx965m 内存条:16G 固态 ...
mongodb集群性能优化
mongodb集群性能优化在前面两篇文章,我们介绍了如何去搭建mongodb集群,这篇文章我们将介绍如何去优化mongodb的各项配置,以达到最优的效果. 警告不做任何的优化,集群搭建完成之后,使 ...
CSV文件保存为utf8编码格式
csv格式文件经常用来批量导入数据到某些应用中,但是经常出现utf8乱码问题,那么该如何解决呢? WPS找不到编码格式设置,微软的office软件有,不过我使用的是libreoffice 步骤如下 1 ...
hostswap dcevm
什么是dcevm dcevm(DynamicCode Evolution Virtual Machine)是java hostspot的补丁(严格上来说是修改),允许(并非无限制)在运行环境下修改加载 ...
求Sn=a+aa+aaa+aaaa+aaaaa的前5项之和，其中a是一个数字
思路:所求和为一个数字的前n项和,例如前4项和就是从4+44+444+4444,一直加到第4位,为4个4.所以可以用一个循环来表示每一项的数字,加到前几项就循环几次.然后将每项进行相加就可以求出总和. ...
免费开源 KiCad EDA 中文资料收集整理（2019-04-30）
免费开源 KiCad EDA 中文资料收集整理用 KiCad 也有一段时间了,为了方便自己查找,整理一下 KiCad 的中文资料,会不定期更新. 会收集KiCad 的新闻.元件封装库.应用技巧.开源 ...
pyspark数据准备
鸢尾花数据集 5.1,3.5,1.4,0.2,Iris-setosa 4.9,3.0,1.4,0.2,Iris-setosa 4.7,3.2,1.3,0.2,Iris-setosa 4.6,3.1,1 ...
使用Jmeter创建ActiveMQ JMS POINT TO POINT请求，环境搭建、请求创建、插件安装、监听服务器资源等
转自:http://www.cnblogs.com/qianyiliushang/p/4348584.html 准备工作: 安装JDK,推荐使用1.7以上版本,并设置JAVA_HOME 下载Jmete ...

Lucene 4.8 - Facet Demo

Lucene 4.8 - Facet Demo的更多相关文章

随机推荐

热门专题