All Downloads are FREE. Search and download functionalities are using the official Maven repository.

com.datatorrent.lib.algo.LeastFrequentKeyMap Maven / Gradle / Ivy

/**
 * Licensed to the Apache Software Foundation (ASF) under one
 * or more contributor license agreements.  See the NOTICE file
 * distributed with this work for additional information
 * regarding copyright ownership.  The ASF licenses this file
 * to you under the Apache License, Version 2.0 (the
 * "License"); you may not use this file except in compliance
 * with the License.  You may obtain a copy of the License at
 *
 *   http://www.apache.org/licenses/LICENSE-2.0
 *
 * Unless required by applicable law or agreed to in writing,
 * software distributed under the License is distributed on an
 * "AS IS" BASIS, WITHOUT WARRANTIES OR CONDITIONS OF ANY
 * KIND, either express or implied.  See the License for the
 * specific language governing permissions and limitations
 * under the License.
 */
package com.datatorrent.lib.algo;

import java.util.ArrayList;
import java.util.HashMap;
import java.util.Map;

import com.datatorrent.api.DefaultInputPort;
import com.datatorrent.api.DefaultOutputPort;
import com.datatorrent.api.annotation.OperatorAnnotation;
import com.datatorrent.api.annotation.OutputPortFieldAnnotation;

import com.datatorrent.lib.util.AbstractBaseFrequentKey;
import com.datatorrent.lib.util.UnifierArrayHashMapFrequent;
import com.datatorrent.lib.util.UnifierHashMapFrequent;

/**
 * This operator filters the incoming stream of key value pairs by finding the key or keys (if there is a tie) that occur the fewest number of times within each window. 
 * A list of the corresponding key value pairs are then output to the port named "list" and one of the corresponding key value pairs is output to the port "least", at the end of each window.
 * 

* Occurrences of each key is counted and at the end of window any of the least frequent key is emitted on output port least and all least frequent * keys on output port list. *

*

* This module is an end of window module. In case of a tie any of the least key would be emitted. The list port would however have all the tied keys
*
* StateFull : Yes, tuple are compared across application window(s).
* Partitions : Yes, least keys are unified on output port.
*
* Ports:
* data: expects Map<K,V>, V is ignored/not used
* least: emits HashMap<K,Integer>(1); where String is the least frequent key, and Integer is the number of its occurrences in the window
* list: emits ArrayList<HashMap<K,Integer>(1)>; Where the list includes all the keys are least frequent
*
*

* * @displayName Emit Least Frequent Tuple Key * @category Rules and Alerts * @tags filter, key value, count * * @since 0.3.2 */ @OperatorAnnotation(partitionable = true) public class LeastFrequentKeyMap extends AbstractBaseFrequentKey { /** * The input port on which key value pairs are received. */ public final transient DefaultInputPort> data = new DefaultInputPort>() { /** * Calls super.processTuple(tuple) for each key in the HashMap */ @Override public void process(Map tuple) { for (Map.Entry e: tuple.entrySet()) { processTuple(e.getKey()); } } }; /** * The output port on which one of the tuples, * which occurred the least number of times, * is emitted. */ @OutputPortFieldAnnotation(optional = true) public final transient DefaultOutputPort> least = new DefaultOutputPort>() { @Override public Unifier> getUnifier() { Unifier> ret = new UnifierHashMapFrequent(); ((UnifierHashMapFrequent)ret).setLeast(true); return ret; } }; /** * The output port on which all the tuples, * which occurred the least number of times, * is emitted. */ @OutputPortFieldAnnotation(optional = true) public final transient DefaultOutputPort>> list = new DefaultOutputPort>>() { @Override public Unifier>> getUnifier() { Unifier>> ret = new UnifierArrayHashMapFrequent(); ((UnifierArrayHashMapFrequent)ret).setLeast(true); return ret; } }; /** * Emits tuple on port "least" * * @param tuple */ @Override public void emitTuple(HashMap tuple) { least.emit(tuple); } /** * Emits tuple on port "list" * * @param tlist */ @Override public void emitList(ArrayList> tlist) { list.emit(tlist); } /** * returns val1 < val2 * * @param val1 * @param val2 * @return val1 < val2 */ @Override public boolean compareCount(int val1, int val2) { return val1 < val2; } }




© 2015 - 2025 Weber Informatics LLC | Privacy Policy