kafka 自定义分区器

package cn.xiaojf.kafka.producer;

import org.apache.kafka.clients.producer.Partitioner;

import org.apache.kafka.common.Cluster;

import org.apache.kafka.common.PartitionInfo;

import org.apache.kafka.common.utils.Utils;

import java.util.List;

import java.util.Map;

import java.util.Random;

import java.util.concurrent.ConcurrentHashMap;

import java.util.concurrent.ConcurrentMap;

import java.util.concurrent.atomic.AtomicInteger;

/**

 * 自定义分区方式

 */

public class CustomPartitioner implements Partitioner {

    private final ConcurrentMap<String, AtomicInteger> topicCounterMap = new ConcurrentHashMap();

    public CustomPartitioner() {

    }

    public void configure(Map<String, ?> configs) {

    }

    /**

     * 自定义分区规则

     * @param topic

     * @param key

     * @param keyBytes

     * @param value

     * @param valueBytes

     * @param cluster

     * @return

     */

    public int partition(String topic, Object key, byte[] keyBytes, Object value, byte[] valueBytes, Cluster cluster) {

        List partitions = cluster.partitionsForTopic(topic);

        int numPartitions = partitions.size();

        if(keyBytes == null) {

            int nextValue = this.nextValue(topic);

            List availablePartitions = cluster.availablePartitionsForTopic(topic);

            if(availablePartitions.size() > 0) {

                int part = Utils.toPositive(nextValue) % availablePartitions.size();

                return ((PartitionInfo)availablePartitions.get(part)).partition();

            } else {

                return Utils.toPositive(nextValue) % numPartitions;

            }

        } else {

            return Utils.toPositive(Utils.murmur2(keyBytes)) % numPartitions;

        }

    }

    private int nextValue(String topic) {

        AtomicInteger counter = (AtomicInteger)this.topicCounterMap.get(topic);

        if(null == counter) {

            counter = new AtomicInteger((new Random()).nextInt());

            AtomicInteger currentCounter = (AtomicInteger)this.topicCounterMap.putIfAbsent(topic, counter);

            if(currentCounter != null) {

                counter = currentCounter;

            }

        }

        return counter.getAndIncrement();

    }

    public void close() {

    }

}

package cn.xiaojf.kafka.producer;

import org.apache.kafka.clients.producer.*;

import org.apache.kafka.common.Cluster;

import org.apache.kafka.common.PartitionInfo;

import org.apache.kafka.common.serialization.StringSerializer;

import org.apache.kafka.common.utils.Utils;

import java.util.*;

import java.util.concurrent.ConcurrentHashMap;

import java.util.concurrent.ConcurrentMap;

import java.util.concurrent.atomic.AtomicInteger;

/**

 * 消息生产者

 * @author xiaojf 2017/3/22 14:27

 */

public class MsgProducer extends Thread {

    private final KafkaProducer<String, String> producer;

    private final String topic;

    private final boolean isAsync;

    public MsgProducer(String topic, boolean isAsync) {

        Properties properties = new Properties();

        properties.put(ProducerConfig.BOOTSTRAP_SERVERS_CONFIG, "192.168.59.130:9092");//broker 集群地址

        properties.put(ProducerConfig.CLIENT_ID_CONFIG, "MsgProducer");//自定义客户端id

        properties.put(ProducerConfig.KEY_SERIALIZER_CLASS_CONFIG, "org.apache.kafka.common.serialization.StringSerializer");//key 序列号方式

        properties.put(ProducerConfig.VALUE_SERIALIZER_CLASS_CONFIG, "org.apache.kafka.common.serialization.StringSerializer");//value 序列号方式

        properties.put(ProducerConfig.PARTITIONER_CLASS_CONFIG,CustomPartitioner.class.getCanonicalName());//自定义分区函数

//        properties.load("properties配置文件");

        this.producer = new KafkaProducer<String, String>(properties);

        this.topic = topic;

        this.isAsync = isAsync;

    }

    @Override

    public void run() {

        int msgNo = 0;

        while (true) {

            String msg = "Msg: " + msgNo;

            String key = msgNo + "";

            if (isAsync) {//异步

                producer.send(new ProducerRecord<String, String>(this.topic,msg));

//                producer.send(new ProducerRecord<String, String>(this.topic, key, msg));

            } else {//同步

                producer.send(new ProducerRecord<String, String>(this.topic, key, msg),

                        new MsgProducerCallback(System.currentTimeMillis(), key, msg));

            }

        }

    }

    /**

     * 消息发送后的回调函数

     */

    class MsgProducerCallback implements Callback {

        private final long startTime;

        private final String key;

        private final String msg;

        public MsgProducerCallback(long startTime, String key, String msg) {

            this.startTime = startTime;

            this.key = key;

            this.msg = msg;

        }

        public void onCompletion(RecordMetadata recordMetadata, Exception e) {

            long elapsedTime = System.currentTimeMillis() - startTime;

            if (recordMetadata != null) {

                System.out.println(msg + " be sended to partition no : " + recordMetadata.partition());

            }

        }

    }

    public static void main(String args[]) {

        new MsgProducer("my-replicated-topic",true).start();//开始发送消息

    }

}

kafka 自定义分区器的更多相关文章

kafka自定义序列化器
<kafka权威指南> Customer.java public class Customer { private int customId; private String custome ...
spark自定义分区器实现
在spark中,框架默认使用的事hashPartitioner分区器进行对rdd分区,但是实际生产中,往往使用spark自带的分区器会产生数据倾斜等原因,这个时候就需要我们自定义分区,按照我们指定的字 ...
MapReduce之自定义分区器Partitioner
@ 目录问题引出默认Partitioner分区自定义Partitioner步骤 Partition分区案例实操分区总结问题引出要求将统计结果按照条件输出到不同文件中(分区). 比如:将统计 ...
玩转Kafka的生产者——分区器与多线程
上篇文章学习kafka的基本安装和基础概念,本文主要是学习kafka的常用API.其中包括生产者和消费者, 多线程生产者,多线程消费者,自定义分区等,当然还包括一些避坑指南. 首发于个人网站:链接地址 ...
kafka producer partitions分区器(七)
消息在经过拦截器.序列化后,就需要确定它发往哪个分区,如果在ProducerRecord中指定了partition字段,那么就不再需要partitioner分区器进行分区了,如果没有指定,那么会根据k ...
Spark源码分析之分区器的作用
最近因为手抖,在Spark中给自己挖了一个数据倾斜的坑.为了解决这个问题,顺便研究了下Spark分区器的原理,趁着周末加班总结一下~ 先说说数据倾斜数据倾斜是指Spark中的RDD在计算的时候,每个 ...
RDD(六)——分区器
RDD的分区器 Spark目前支持Hash分区和Range分区,用户也可以自定义分区,Hash分区为当前的默认分区,Spark中分区器直接决定了RDD中分区的个数.RDD中每条数据经过Shuffle过 ...
Kafka的接口回调 +自定义分区、拦截器
一.接口回调+自定义分区 1.接口回调:在使用消费者的send方法时添加Callback回调 producer.send(new ProducerRecord<String, String> ...
【Kafka】自定义分区策略
自定义分区策略思路 Command+Option+shift+N 调出查询页面,找到producer包的Partitioner接口 Partitioner下有一个DefaultPartitioner ...

随机推荐

这个demo是为解决IQKeyboardManager和Masonry同时使用时，导航栏上移和make.right失效的问题
原文链接在我的个人博客主页 (一).引言: 在 IQKeyboardManager 和 Masonry 同时使用时,导航栏上移和make.right失效等问题多多. 其实我们完美的效果应该是这样的:* ...
图解函数重载以及arguments
Android -- Annotation（注解）原理详解及常见框架应用
1,我们在上一篇讲到了EventBus源码及3.0版本的简单使用,知道了我们3.0版本是使用注解方式标记事件响应方法的,这里我们就有一个疑问了,为什么在一个方法加上类似于"@Subscrib ...
AFNetworking 内部详解
AFNetworking 是一个适用于IOS 和 Mac OSX 两个平台的网络库,他是在Foundation URL Loading System 基础上进行的一套封装 ,并提供了丰富的API接口 ...
Java 原始数据类型的计算：运算符重载（Operator Overload）和类型转换（Type Conversion)
原文阅读:<算法(第四版)>第一章第一节:基础编程模型有没有在面试的时候被问到:下面这几行代码的执行结果是什么?依据是什么? System.out.println (5/3); Sys ...
Mvc自定义验证
假设我们书店需要录入一本书,为了简单的体现我们的自定义验证,我们的实体定义的非常简单,就两个属性,一个名称Name,一个出版社Publisher. public class BookInfo { pu ...
读书笔记 effective c++ Item 51 实现new和delete的时候要遵守约定
Item 50中解释了在什么情况下你可能想实现自己版本的operator new和operator delete,但是没有解释当你实现的时候需要遵守的约定.遵守这些规则并不是很困难,但是它们其中有一些 ...
python与opencv的结合之人脸识别值
首先还是要感谢http://www.jb51.net/article/67392.htm这位大神的无私奉献,开源的代码,让我省去了很多事,但是就光系统环境的配置就花去了我将近一个星期的时间,真是不容易 ...
java类集框架(ArrayList,LinkedList,Vector区别)
主要分两个接口:collection和Map 主要分三类:集合(set).列表(List).映射(Map)1.集合:没有重复对象,没有特定排序方式2.列表:对象按索引位置排序,可以有重复对象3.映射: ...
SQLite数据库_实现简单的增删改查
1.SQLite是一款轻量型的数据库是遵守ACID(原子性.一致性.隔离性.持久性)的关联式数据库管理系统,多用于嵌入式开发中. 2.Android平台中嵌入了一个关系型数据库SQLite,和其他数据 ...

kafka 自定义分区器

kafka 自定义分区器的更多相关文章

随机推荐

热门专题