Python SimilarStreamAggregation.SimilarStreamAggregation示例

编程语言: Python

命名空间/包名称: experiments.ssa.ssa

方法/功能: SimilarStreamAggregation

hotexamples.com的示例: 4

Python SimilarStreamAggregation.SimilarStreamAggregation - 已找到4个示例。这些是从开源项目中提取的最受好评的experiments.ssa.ssa.SimilarStreamAggregation.SimilarStreamAggregation现实Python示例。您可以评价示例，以帮助我们提高示例质量。

常用方法

显示隐藏

SimilarStreamAggregation(4)

estimate(4)

iterateClusters(3)

示例#1

显示文件

文件： quality_comparison_with_ssa.py 项目： ylaron/hd_streams_clustering

 def getStatsForSSA(self):
     print 'SSA'
     ts = time.time()
     sstObject = SimilarStreamAggregation(
         dict(self._iterateUserDocuments()),
         self.stream_settings['ssa_threshold'])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return self.getEvaluationMetrics(documentClusters, te - ts)

示例#2

显示文件

 def performanceForCDAITAt(noOfTweets, fileName, **stream_settings):
     ts = time.time()
     sstObject = SimilarStreamAggregation(
         dict(
             iterateTweetUsersAfterCombiningTweets(fileName,
                                                   **stream_settings)),
         stream_settings['ssa_threshold'])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return Evaluation.getEvaluationMetrics(noOfTweets, documentClusters,
                                            te - ts)

示例#3

显示文件

def getStatsForSSA():
    batchSize = 10000
    default_experts_twitter_stream_settings['ssa_threshold'] = 0.75
    for id in range(21, 50):
        fileName = time_to_process_points + '%s/%s' % (batchSize, id)
        ts = time.time()
        sstObject = SimilarStreamAggregation(
            dict(iterateUserDocuments(fileName)),
            default_experts_twitter_stream_settings['ssa_threshold'])
        sstObject.estimate()
        #    documentClusters = list(sstObject.iterateClusters())
        iteration_data = {
            'iteration_time': time.time() - ts,
            'type': 'ssa',
            'number_of_messages': batchSize * (id + 1),
            'batch_size': batchSize
        }
        FileIO.writeToFileAsJson(iteration_data, ssa_stats_file)

示例#4

显示文件

文件： ssa_tests.py 项目： ylaron/hd_streams_clustering

 def test_estimate(self):
     nn = SimilarStreamAggregation(vectors, 0.99)
     nn.estimate()
     self.assertEqual([['1', '3', '2'], ['5', '7']], list(nn.iterateClusters()))