Python SimilarStreamAggregation.estimate示例

编程语言: Python

命名空间/包名称: experiments.ssa.ssa

方法/功能: estimate

hotexamples.com的示例: 7

Python SimilarStreamAggregation.estimate - 已找到7个示例。这些是从开源项目中提取的最受好评的experiments.ssa.ssa.SimilarStreamAggregation.estimate现实Python示例。您可以评价示例，以帮助我们提高示例质量。

常用方法

显示隐藏

SimilarStreamAggregation(4)

estimate(4)

iterateClusters(3)

示例#1

显示文件

文件： performance_with_cda.py 项目： greeness/hd_streams_clustering

 def performanceForCDAITAt(noOfTweets, fileName, **stream_settings):
     ts = time.time()
     sstObject = SimilarStreamAggregation(dict(iterateTweetUsersAfterCombiningTweets(fileName, **stream_settings)), stream_settings['ssa_threshold'])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return Evaluation.getEvaluationMetrics(noOfTweets, documentClusters, te-ts)

示例#2

显示文件

文件： quality_comparison_with_ssa.py 项目： kykamath/hd_streams_clustering

 def getStatsForSSA(self):
     print "SSA"
     ts = time.time()
     sstObject = SimilarStreamAggregation(dict(self._iterateUserDocuments()), self.stream_settings["ssa_threshold"])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return self.getEvaluationMetrics(documentClusters, te - ts)

示例#3

显示文件

文件： quality_comparison_with_ssa.py 项目： ylaron/hd_streams_clustering

 def getStatsForSSA(self):
     print 'SSA'
     ts = time.time()
     sstObject = SimilarStreamAggregation(
         dict(self._iterateUserDocuments()),
         self.stream_settings['ssa_threshold'])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return self.getEvaluationMetrics(documentClusters, te - ts)

示例#4

显示文件

文件： time_to_process_points.py 项目： kykamath/hd_streams_clustering

def getStatsForSSA():
    batchSize = 10000
    default_experts_twitter_stream_settings['ssa_threshold']=0.75
    for id in range(21, 50):
        fileName = time_to_process_points+'%s/%s'%(batchSize,id)
        ts = time.time()
        sstObject = SimilarStreamAggregation(dict(iterateUserDocuments(fileName)), default_experts_twitter_stream_settings['ssa_threshold'])
        sstObject.estimate()
    #    documentClusters = list(sstObject.iterateClusters())
        iteration_data = {'iteration_time': time.time()-ts, 'type': 'ssa', 'number_of_messages': batchSize*(id+1), 'batch_size': batchSize}
        FileIO.writeToFileAsJson(iteration_data, ssa_stats_file)

示例#5

显示文件

 def performanceForCDAITAt(noOfTweets, fileName, **stream_settings):
     ts = time.time()
     sstObject = SimilarStreamAggregation(
         dict(
             iterateTweetUsersAfterCombiningTweets(fileName,
                                                   **stream_settings)),
         stream_settings['ssa_threshold'])
     sstObject.estimate()
     documentClusters = list(sstObject.iterateClusters())
     te = time.time()
     return Evaluation.getEvaluationMetrics(noOfTweets, documentClusters,
                                            te - ts)

示例#6

显示文件

def getStatsForSSA():
    batchSize = 10000
    default_experts_twitter_stream_settings['ssa_threshold'] = 0.75
    for id in range(21, 50):
        fileName = time_to_process_points + '%s/%s' % (batchSize, id)
        ts = time.time()
        sstObject = SimilarStreamAggregation(
            dict(iterateUserDocuments(fileName)),
            default_experts_twitter_stream_settings['ssa_threshold'])
        sstObject.estimate()
        #    documentClusters = list(sstObject.iterateClusters())
        iteration_data = {
            'iteration_time': time.time() - ts,
            'type': 'ssa',
            'number_of_messages': batchSize * (id + 1),
            'batch_size': batchSize
        }
        FileIO.writeToFileAsJson(iteration_data, ssa_stats_file)

示例#7

显示文件

文件： ssa_tests.py 项目： ylaron/hd_streams_clustering

 def test_estimate(self):
     nn = SimilarStreamAggregation(vectors, 0.99)
     nn.estimate()
     self.assertEqual([['1', '3', '2'], ['5', '7']], list(nn.iterateClusters()))